Merge pull request #1767 from xmy0916/dygraph

add multi language config file imgs and dict
4 years ago · 21292bb3b7
parent 301b0d347c cafea5dcbe
commit 21292bb3b7
68 changed files with 11359 additions and 29 deletions
--- a/configs/rec/multi_language/generate_multi_language_configs.py
+++ b/configs/rec/multi_language/generate_multi_language_configs.py
@ -0,0 +1,152 @@
+# Copyright (c) 2021 PaddlePaddle Authors. All Rights Reserved.
+#
+# Licensed under the Apache License, Version 2.0 (the "License");
+# you may not use this file except in compliance with the License.
+# You may obtain a copy of the License at
+#
+#     http://www.apache.org/licenses/LICENSE-2.0
+#
+# Unless required by applicable law or agreed to in writing, software
+# distributed under the License is distributed on an "AS IS" BASIS,
+# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
+# See the License for the specific language governing permissions and
+# limitations under the License.
+
+import yaml
+from argparse import ArgumentParser, RawDescriptionHelpFormatter
+import os.path
+import logging
+logging.basicConfig(level=logging.INFO)
+
+support_list = {
+    'it':'italian', 'xi':'spanish', 'pu':'portuguese', 'ru':'russian', 'ar':'arabic',
+    'ta':'tamil', 'ug':'uyghur', 'fa':'persian', 'ur':'urdu', 'rs':'serbian latin',
+    'oc':'occitan', 'rsc':'serbian cyrillic', 'bg':'bulgarian', 'uk':'ukranian', 'be':'belarusian',
+    'te':'telugu', 'ka':'kannada', 'chinese_cht':'chinese tradition','hi':'hindi','mr':'marathi',
+    'ne':'nepali',
+}
+assert(
+    os.path.isfile("./rec_multi_language_lite_train.yml")
+    ),"Loss basic configuration file rec_multi_language_lite_train.yml.\
+You can download it from \
+https://github.com/PaddlePaddle/PaddleOCR/tree/dygraph/configs/rec/multi_language/"
+ 
+global_config = yaml.load(open("./rec_multi_language_lite_train.yml", 'rb'), Loader=yaml.Loader)
+project_path = os.path.abspath(os.path.join(os.getcwd(), "../../../"))
+
+class ArgsParser(ArgumentParser):
+    def __init__(self):
+        super(ArgsParser, self).__init__(
+            formatter_class=RawDescriptionHelpFormatter)
+        self.add_argument(
+            "-o", "--opt", nargs='+', help="set configuration options")
+        self.add_argument(
+            "-l", "--language", nargs='+', help="set language type, support {}".format(support_list))
+        self.add_argument(
+            "--train",type=str,help="you can use this command to change the train dataset default path")
+        self.add_argument(
+            "--val",type=str,help="you can use this command to change the eval dataset default path")
+        self.add_argument(
+            "--dict",type=str,help="you can use this command to change the dictionary default path")
+        self.add_argument(
+            "--data_dir",type=str,help="you can use this command to change the dataset default root path")
+
+    def parse_args(self, argv=None):
+        args = super(ArgsParser, self).parse_args(argv)
+        args.opt = self._parse_opt(args.opt)
+        args.language = self._set_language(args.language)
+        return args
+
+    def _parse_opt(self, opts):
+        config = {}
+        if not opts:
+            return config
+        for s in opts:
+            s = s.strip()
+            k, v = s.split('=')
+            config[k] = yaml.load(v, Loader=yaml.Loader)
+        return config
+
+    def _set_language(self, type):
+        assert(type),"please use -l or --language to choose language type"
+        assert(
+                type[0] in support_list.keys()
+               ),"the sub_keys(-l or --language) can only be one of support list: \n{},\nbut get: {}, " \
+                 "please check your running command".format(support_list, type)
+        global_config['Global']['character_dict_path'] = 'ppocr/utils/dict/{}_dict.txt'.format(type[0])
+        global_config['Global']['save_model_dir'] = './output/rec_{}_lite'.format(type[0])
+        global_config['Train']['dataset']['label_file_list'] = ["train_data/{}_train.txt".format(type[0])]
+        global_config['Eval']['dataset']['label_file_list'] = ["train_data/{}_val.txt".format(type[0])]
+        global_config['Global']['character_type'] = type[0]
+        assert(
+                os.path.isfile(os.path.join(project_path,global_config['Global']['character_dict_path']))
+              ),"Loss default dictionary file {}_dict.txt.You can download it from \
+https://github.com/PaddlePaddle/PaddleOCR/tree/dygraph/ppocr/utils/dict/".format(type[0])
+        return type[0]
+
+
+def merge_config(config):
+    """
+    Merge config into global config.
+    Args:
+        config (dict): Config to be merged.
+    Returns: global config
+    """
+    for key, value in config.items():
+        if "." not in key:
+            if isinstance(value, dict) and key in global_config:
+                global_config[key].update(value)
+            else:
+                global_config[key] = value
+        else:
+            sub_keys = key.split('.')
+            assert (
+                sub_keys[0] in global_config
+            ), "the sub_keys can only be one of global_config: {}, but get: {}, please check your running command".format(
+                global_config.keys(), sub_keys[0])
+            cur = global_config[sub_keys[0]]
+            for idx, sub_key in enumerate(sub_keys[1:]):
+                if idx == len(sub_keys) - 2:
+                    cur[sub_key] = value
+                else:
+                    cur = cur[sub_key]
+                    
+def loss_file(path):
+    assert(
+            os.path.exists(path)
+          ),"There is no such file:{},Please do not forget to put in the specified file".format(path)
+
+        
+if __name__ == '__main__':
+    FLAGS = ArgsParser().parse_args()
+    merge_config(FLAGS.opt)
+    save_file_path = 'rec_{}_lite_train.yml'.format(FLAGS.language)
+    if os.path.isfile(save_file_path):
+        os.remove(save_file_path)
+        
+    if FLAGS.train:
+        global_config['Train']['dataset']['label_file_list'] = [FLAGS.train]
+        train_label_path = os.path.join(project_path,FLAGS.train)
+        loss_file(train_label_path)
+    if FLAGS.val:
+        global_config['Eval']['dataset']['label_file_list'] = [FLAGS.val]
+        eval_label_path = os.path.join(project_path,FLAGS.val)
+        loss_file(Eval_label_path)
+    if FLAGS.dict:
+        global_config['Global']['character_dict_path'] = FLAGS.dict
+        dict_path = os.path.join(project_path,FLAGS.dict)
+        loss_file(dict_path)
+    if FLAGS.data_dir:
+        global_config['Eval']['dataset']['data_dir'] = FLAGS.data_dir
+        global_config['Train']['dataset']['data_dir'] = FLAGS.data_dir
+        data_dir = os.path.join(project_path,FLAGS.data_dir)
+        loss_file(data_dir)
+        
+    with open(save_file_path, 'w') as f:
+        yaml.dump(dict(global_config), f, default_flow_style=False, sort_keys=False)
+    logging.info("Project path is          :{}".format(project_path))
+    logging.info("Train list path set to   :{}".format(global_config['Train']['dataset']['label_file_list'][0]))
+    logging.info("Eval list path set to    :{}".format(global_config['Eval']['dataset']['label_file_list'][0]))
+    logging.info("Dataset root path set to :{}".format(global_config['Eval']['dataset']['data_dir']))
+    logging.info("Dict path set to         :{}".format(global_config['Global']['character_dict_path']))
+    logging.info("Config file set to       :configs/rec/multi_language/{}".format(save_file_path))
--- a/configs/rec/multi_language/rec_multi_language_lite_train.yml
+++ b/configs/rec/multi_language/rec_multi_language_lite_train.yml
@ -0,0 +1,103 @@
+Global:
+  use_gpu: True
+  epoch_num: 500
+  log_smooth_window: 20
+  print_batch_step: 10
+  save_model_dir: ./output/rec_multi_language_lite
+  save_epoch_step: 3
+  # evaluation is run every 5000 iterations after the 4000th iteration
+  eval_batch_step: [0, 2000]
+  # if pretrained_model is saved in static mode, load_static_weights must set to True
+  cal_metric_during_train: True
+  pretrained_model: 
+  checkpoints: 
+  save_inference_dir: 
+  use_visualdl: False
+  infer_img:
+  # for data or label process
+  character_dict_path: 
+  # Set the language of training, if set, select the default dictionary file
+  character_type: 
+  max_text_length: 25
+  infer_mode: False
+  use_space_char: True
+
+
+Optimizer:
+  name: Adam
+  beta1: 0.9
+  beta2: 0.999
+  lr:
+    name: Cosine
+    learning_rate: 0.001
+  regularizer:
+    name: 'L2'
+    factor: 0.00001
+
+Architecture:
+  model_type: rec
+  algorithm: CRNN
+  Transform:
+  Backbone:
+    name: MobileNetV3
+    scale: 0.5
+    model_name: small
+    small_stride: [1, 2, 2, 2]
+  Neck:
+    name: SequenceEncoder
+    encoder_type: rnn
+    hidden_size: 48
+  Head:
+    name: CTCHead
+    fc_decay: 0.00001
+
+Loss:
+  name: CTCLoss
+
+PostProcess:
+  name: CTCLabelDecode
+
+Metric:
+  name: RecMetric
+  main_indicator: acc
+
+Train:
+  dataset:
+    name: SimpleDataSet
+    data_dir: train_data/
+    label_file_list: ["./train_data/train_list.txt"]
+    transforms:
+      - DecodeImage: # load image
+          img_mode: BGR
+          channel_first: False
+      - RecAug: 
+      - CTCLabelEncode: # Class handling label
+      - RecResizeImg:
+          image_shape: [3, 32, 320]
+      - KeepKeys:
+          keep_keys: ['image', 'label', 'length'] # dataloader will return list in this order
+  loader:
+    shuffle: True
+    batch_size_per_card: 256
+    drop_last: True
+    num_workers: 8
+
+Eval:
+  dataset:
+    name: SimpleDataSet
+    data_dir: train_data/
+    label_file_list: ["./train_data/val_list.txt"]
+    transforms:
+      - DecodeImage: # load image
+          img_mode: BGR
+          channel_first: False
+      - CTCLabelEncode: # Class handling label
+      - RecResizeImg:
+          image_shape: [3, 32, 320]
+      - KeepKeys:
+          keep_keys: ['image', 'label', 'length'] # dataloader will return list in this order
+  loader:
+    shuffle: False
+    drop_last: False
+    batch_size_per_card: 256
+    num_workers: 8
--- a/doc/doc_ch/models_list.md
+++ b/doc/doc_ch/models_list.md
@ -60,6 +60,7 @@ PaddleOCR提供的可下载模型包括`推理模型`、`训练模型`、`预训
 | japan_mobile_v2.0_rec |日文识别|[rec_japan_lite_train.yml](../../configs/rec/multi_language/rec_japan_lite_train.yml)|4.23M|[推理模型](https://paddleocr.bj.bcebos.com/dygraph_v2.0/multilingual/japan_mobile_v2.0_rec_infer.tar) / [训练模型](https://paddleocr.bj.bcebos.com/dygraph_v2.0/multilingual/japan_mobile_v2.0_rec_train.tar) |


+
 <a name="文本方向分类模型"></a>
 ### 三、文本方向分类模型

--- a/doc/imgs/arabic_1.jpg
+++ b/doc/imgs/arabic_1.jpg
--- a/doc/imgs/arabic_2.jpg
+++ b/doc/imgs/arabic_2.jpg
--- a/doc/imgs/be_1.jpg
+++ b/doc/imgs/be_1.jpg
--- a/doc/imgs/be_2.jpg
+++ b/doc/imgs/be_2.jpg
--- a/doc/imgs/bg_1.jpg
+++ b/doc/imgs/bg_1.jpg
--- a/doc/imgs/bg_2.jpg
+++ b/doc/imgs/bg_2.jpg
--- a/doc/imgs/chinese_cht_1.png
+++ b/doc/imgs/chinese_cht_1.png
--- a/doc/imgs/chinese_cht_2.png
+++ b/doc/imgs/chinese_cht_2.png
--- a/doc/imgs/fa_1.jpg
+++ b/doc/imgs/fa_1.jpg
--- a/doc/imgs/fa_2.jpg
+++ b/doc/imgs/fa_2.jpg
--- a/doc/imgs/hi_1.jpg
+++ b/doc/imgs/hi_1.jpg
--- a/doc/imgs/hi_2.jpg
+++ b/doc/imgs/hi_2.jpg
--- a/doc/imgs/it_1.jpg
+++ b/doc/imgs/it_1.jpg
--- a/doc/imgs/it_2.jpg
+++ b/doc/imgs/it_2.jpg
--- a/doc/imgs/ka_1.jpg
+++ b/doc/imgs/ka_1.jpg
--- a/doc/imgs/ka_2.jpg
+++ b/doc/imgs/ka_2.jpg
--- a/doc/imgs/mr_1.jpg
+++ b/doc/imgs/mr_1.jpg
--- a/doc/imgs/mr_2.jpg
+++ b/doc/imgs/mr_2.jpg
--- a/doc/imgs/ne_1.jpg
+++ b/doc/imgs/ne_1.jpg
--- a/doc/imgs/ne_2.jpg
+++ b/doc/imgs/ne_2.jpg
--- a/doc/imgs/oc_1.jpg
+++ b/doc/imgs/oc_1.jpg
--- a/doc/imgs/oc_2.jpg
+++ b/doc/imgs/oc_2.jpg
--- a/Show More
+++ b/Show More