Merge remote-tracking branch 'upstream/dygraph' into dy1

PaddlePaddle · Jan 25, 2021 · 6fb6152 · 6fb6152
2 parents 1fd7e11 + a78afa2
commit 6fb6152
Show file tree

Hide file tree

Showing 84 changed files with 11,782 additions and 226 deletions.
diff --git a/README_ch.md b/README_ch.md
@@ -8,7 +8,7 @@ PaddleOCR同时支持动态图与静态图两种编程范式
 - 静态图版本：develop分支
 
 **近期更新**
-- 2021.1.11 [FAQ](./doc/doc_ch/FAQ.md)新增5个高频问题，总数147个，每周一都会更新，欢迎大家持续关注。
+- 2021.1.18 [FAQ](./doc/doc_ch/FAQ.md)新增5个高频问题，总数152个，每周一都会更新，欢迎大家持续关注。
 - 2020.12.15 更新数据合成工具[Style-Text](./StyleText/README_ch.md)，可以批量合成大量与目标场景类似的图像，在多个场景验证，效果明显提升。
 - 2020.11.25 更新半自动标注工具[PPOCRLabel](./PPOCRLabel/README_ch.md)，辅助开发者高效完成标注任务，输出格式与PP-OCR训练任务完美衔接。
 - 2020.9.22 更新PP-OCR技术文章，https://arxiv.org/abs/2009.09941
@@ -101,8 +101,8 @@ PaddleOCR同时支持动态图与静态图两种编程范式
 - [效果展示](#效果展示)
 - FAQ
  - [【精选】OCR精选10个问题](./doc/doc_ch/FAQ.md)
- - [【理论篇】OCR通用31个问题](./doc/doc_ch/FAQ.md)
- - [【实战篇】PaddleOCR实战106个问题](./doc/doc_ch/FAQ.md)
+ - [【理论篇】OCR通用32个问题](./doc/doc_ch/FAQ.md)
+ - [【实战篇】PaddleOCR实战110个问题](./doc/doc_ch/FAQ.md)
 - [技术交流群](#欢迎加入PaddleOCR技术交流群)
 - [参考文献](./doc/doc_ch/reference.md)
 - [许可证书](#许可证书)

diff --git a/configs/rec/multi_language/generate_multi_language_configs.py b/configs/rec/multi_language/generate_multi_language_configs.py
@@ -0,0 +1,152 @@
+# Copyright (c) 2021 PaddlePaddle Authors. All Rights Reserved.
+#
+# Licensed under the Apache License, Version 2.0 (the "License");
+# you may not use this file except in compliance with the License.
+# You may obtain a copy of the License at
+#
+# http:https://www.apache.org/licenses/LICENSE-2.0
+#
+# Unless required by applicable law or agreed to in writing, software
+# distributed under the License is distributed on an "AS IS" BASIS,
+# WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
+# See the License for the specific language governing permissions and
+# limitations under the License.
+
+import yaml
+from argparse import ArgumentParser, RawDescriptionHelpFormatter
+import os.path
+import logging
+logging.basicConfig(level=logging.INFO)
+
+support_list = {
+ 'it':'italian', 'xi':'spanish', 'pu':'portuguese', 'ru':'russian', 'ar':'arabic',
+ 'ta':'tamil', 'ug':'uyghur', 'fa':'persian', 'ur':'urdu', 'rs':'serbian latin',
+ 'oc':'occitan', 'rsc':'serbian cyrillic', 'bg':'bulgarian', 'uk':'ukranian', 'be':'belarusian',
+ 'te':'telugu', 'ka':'kannada', 'chinese_cht':'chinese tradition','hi':'hindi','mr':'marathi',
+ 'ne':'nepali',
+}
+assert(
+ os.path.isfile("./rec_multi_language_lite_train.yml")
+ ),"Loss basic configuration file rec_multi_language_lite_train.yml.\
+You can download it from \
+https://github.com/PaddlePaddle/PaddleOCR/tree/dygraph/configs/rec/multi_language/"
+
+global_config = yaml.load(open("./rec_multi_language_lite_train.yml", 'rb'), Loader=yaml.Loader)
+project_path = os.path.abspath(os.path.join(os.getcwd(), "../../../"))
+
+class ArgsParser(ArgumentParser):
+ def __init__(self):
+ super(ArgsParser, self).__init__(
+ formatter_class=RawDescriptionHelpFormatter)
+ self.add_argument(
+ "-o", "--opt", nargs='+', help="set configuration options")
+ self.add_argument(
+ "-l", "--language", nargs='+', help="set language type, support {}".format(support_list))
+ self.add_argument(
+ "--train",type=str,help="you can use this command to change the train dataset default path")
+ self.add_argument(
+ "--val",type=str,help="you can use this command to change the eval dataset default path")
+ self.add_argument(
+ "--dict",type=str,help="you can use this command to change the dictionary default path")
+ self.add_argument(
+ "--data_dir",type=str,help="you can use this command to change the dataset default root path")
+
+ def parse_args(self, argv=None):
+ args = super(ArgsParser, self).parse_args(argv)
+ args.opt = self._parse_opt(args.opt)
+ args.language = self._set_language(args.language)
+ return args
+
+ def _parse_opt(self, opts):
+ config = {}
+ if not opts:
+ return config
+ for s in opts:
+ s = s.strip()
+ k, v = s.split('=')
+ config[k] = yaml.load(v, Loader=yaml.Loader)
+ return config
+
+ def _set_language(self, type):
+ assert(type),"please use -l or --language to choose language type"
+ assert(
+ type[0] in support_list.keys()
+ ),"the sub_keys(-l or --language) can only be one of support list: \n{},\nbut get: {}, " \
+ "please check your running command".format(support_list, type)
+ global_config['Global']['character_dict_path'] = 'ppocr/utils/dict/{}_dict.txt'.format(type[0])
+ global_config['Global']['save_model_dir'] = './output/rec_{}_lite'.format(type[0])
+ global_config['Train']['dataset']['label_file_list'] = ["train_data/{}_train.txt".format(type[0])]
+ global_config['Eval']['dataset']['label_file_list'] = ["train_data/{}_val.txt".format(type[0])]
+ global_config['Global']['character_type'] = type[0]
+ assert(
+ os.path.isfile(os.path.join(project_path,global_config['Global']['character_dict_path']))
+ ),"Loss default dictionary file {}_dict.txt.You can download it from \
+https://github.com/PaddlePaddle/PaddleOCR/tree/dygraph/ppocr/utils/dict/".format(type[0])
+ return type[0]
+
+
+def merge_config(config):
+ """
+ Merge config into global config.
+ Args:
+ config (dict): Config to be merged.
+ Returns: global config
+ """
+ for key, value in config.items():
+ if "." not in key:
+ if isinstance(value, dict) and key in global_config:
+ global_config[key].update(value)
+ else:
+ global_config[key] = value
+ else:
+ sub_keys = key.split('.')
+ assert (
+ sub_keys[0] in global_config
+ ), "the sub_keys can only be one of global_config: {}, but get: {}, please check your running command".format(
+ global_config.keys(), sub_keys[0])
+ cur = global_config[sub_keys[0]]
+ for idx, sub_key in enumerate(sub_keys[1:]):
+ if idx == len(sub_keys) - 2:
+ cur[sub_key] = value
+ else:
+ cur = cur[sub_key]
+
+def loss_file(path):
+ assert(
+ os.path.exists(path)
+ ),"There is no such file:{},Please do not forget to put in the specified file".format(path)
+
+
+if __name__ == '__main__':
+ FLAGS = ArgsParser().parse_args()
+ merge_config(FLAGS.opt)
+ save_file_path = 'rec_{}_lite_train.yml'.format(FLAGS.language)
+ if os.path.isfile(save_file_path):
+ os.remove(save_file_path)
+
+ if FLAGS.train:
+ global_config['Train']['dataset']['label_file_list'] = [FLAGS.train]
+ train_label_path = os.path.join(project_path,FLAGS.train)
+ loss_file(train_label_path)
+ if FLAGS.val:
+ global_config['Eval']['dataset']['label_file_list'] = [FLAGS.val]
+ eval_label_path = os.path.join(project_path,FLAGS.val)
+ loss_file(Eval_label_path)
+ if FLAGS.dict:
+ global_config['Global']['character_dict_path'] = FLAGS.dict
+ dict_path = os.path.join(project_path,FLAGS.dict)
+ loss_file(dict_path)
+ if FLAGS.data_dir:
+ global_config['Eval']['dataset']['data_dir'] = FLAGS.data_dir
+ global_config['Train']['dataset']['data_dir'] = FLAGS.data_dir
+ data_dir = os.path.join(project_path,FLAGS.data_dir)
+ loss_file(data_dir)
+
+ with open(save_file_path, 'w') as f:
+ yaml.dump(dict(global_config), f, default_flow_style=False, sort_keys=False)
+ logging.info("Project path is :{}".format(project_path))
+ logging.info("Train list path set to :{}".format(global_config['Train']['dataset']['label_file_list'][0]))
+ logging.info("Eval list path set to :{}".format(global_config['Eval']['dataset']['label_file_list'][0]))
+ logging.info("Dataset root path set to :{}".format(global_config['Eval']['dataset']['data_dir']))
+ logging.info("Dict path set to :{}".format(global_config['Global']['character_dict_path']))
+ logging.info("Config file set to :configs/rec/multi_language/{}".format(save_file_path))
diff --git a/configs/rec/multi_language/rec_multi_language_lite_train.yml b/configs/rec/multi_language/rec_multi_language_lite_train.yml
@@ -0,0 +1,103 @@
+Global:
+ use_gpu: True
+ epoch_num: 500
+ log_smooth_window: 20
+ print_batch_step: 10
+ save_model_dir: ./output/rec_multi_language_lite
+ save_epoch_step: 3
+ # evaluation is run every 5000 iterations after the 4000th iteration
+ eval_batch_step: [0, 2000]
+ # if pretrained_model is saved in static mode, load_static_weights must set to True
+ cal_metric_during_train: True
+ pretrained_model: 
+ checkpoints: 
+ save_inference_dir: 
+ use_visualdl: False
+ infer_img:
+ # for data or label process
+ character_dict_path: 
+ # Set the language of training, if set, select the default dictionary file
+ character_type: 
+ max_text_length: 25
+ infer_mode: False
+ use_space_char: True
+
+
+Optimizer:
+ name: Adam
+ beta1: 0.9
+ beta2: 0.999
+ lr:
+ name: Cosine
+ learning_rate: 0.001
+ regularizer:
+ name: 'L2'
+ factor: 0.00001
+
+Architecture:
+ model_type: rec
+ algorithm: CRNN
+ Transform:
+ Backbone:
+ name: MobileNetV3
+ scale: 0.5
+ model_name: small
+ small_stride: [1, 2, 2, 2]
+ Neck:
+ name: SequenceEncoder
+ encoder_type: rnn
+ hidden_size: 48
+ Head:
+ name: CTCHead
+ fc_decay: 0.00001
+
+Loss:
+ name: CTCLoss
+
+PostProcess:
+ name: CTCLabelDecode
+
+Metric:
+ name: RecMetric
+ main_indicator: acc
+
+Train:
+ dataset:
+ name: SimpleDataSet
+ data_dir: train_data/
+ label_file_list: ["./train_data/train_list.txt"]
+ transforms:
+ - DecodeImage: # load image
+ img_mode: BGR
+ channel_first: False
+ - RecAug: 
+ - CTCLabelEncode: # Class handling label
+ - RecResizeImg:
+ image_shape: [3, 32, 320]
+ - KeepKeys:
+ keep_keys: ['image', 'label', 'length'] # dataloader will return list in this order
+ loader:
+ shuffle: True
+ batch_size_per_card: 256
+ drop_last: True
+ num_workers: 8
+
+Eval:
+ dataset:
+ name: SimpleDataSet
+ data_dir: train_data/
+ label_file_list: ["./train_data/val_list.txt"]
+ transforms:
+ - DecodeImage: # load image
+ img_mode: BGR
+ channel_first: False
+ - CTCLabelEncode: # Class handling label
+ - RecResizeImg:
+ image_shape: [3, 32, 320]
+ - KeepKeys:
+ keep_keys: ['image', 'label', 'length'] # dataloader will return list in this order
+ loader:
+ shuffle: False
+ drop_last: False
+ batch_size_per_card: 256
+ num_workers: 8
diff --git a/deploy/docker/hubserving/cpu/Dockerfile b/deploy/docker/hubserving/cpu/Dockerfile
@@ -1,11 +1,9 @@
-# Version: 1.0.0
-FROM hub.baidubce.com/paddlepaddle/paddle:latest-gpu-cuda10.0-cudnn7-dev
+# Version: 2.0.0
+FROM registry.baidubce.com/paddlepaddle/paddle:2.0.0rc1
 
 # PaddleOCR base on Python3.7
 RUN pip3.7 install --upgrade pip -i https://mirror.baidu.com/pypi/simple
 
-RUN python3.7 -m pip install paddlepaddle==2.0.0rc0 -i https://mirror.baidu.com/pypi/simple
-
 RUN pip3.7 install paddlehub --upgrade -i https://mirror.baidu.com/pypi/simple
 
 RUN git clone https://github.com/PaddlePaddle/PaddleOCR.git /PaddleOCR
@@ -15,15 +13,15 @@ WORKDIR /PaddleOCR
 RUN pip3.7 install -r requirements.txt -i https://mirror.baidu.com/pypi/simple
 
 RUN mkdir -p /PaddleOCR/inference/
-# Download orc detect model(light version). if you want to change normal version, you can change ch_ppocr_mobile_v1.1_det_infer to ch_ppocr_server_v1.1_det_infer, also remember change det_model_dir in deploy/hubserving/ocr_system/params.py）
+# Download orc detect model(light version). if you want to change normal version, you can change ch_ppocr_mobile_v2.0_det_infer to ch_ppocr_server_v2.0_det_infer, also remember change det_model_dir in deploy/hubserving/ocr_system/params.py）
 ADD {link} /PaddleOCR/inference/
 RUN tar xf /PaddleOCR/inference/{file} -C /PaddleOCR/inference/
 
-# Download direction classifier(light version). If you want to change normal version, you can change ch_ppocr_mobile_v1.1_cls_infer to ch_ppocr_mobile_v1.1_cls_infer, also remember change cls_model_dir in deploy/hubserving/ocr_system/params.py）
+# Download direction classifier(light version). If you want to change normal version, you can change ch_ppocr_mobile_v2.0_cls_infer to ch_ppocr_mobile_v2.0_cls_infer, also remember change cls_model_dir in deploy/hubserving/ocr_system/params.py）
 ADD {link} /PaddleOCR/inference/
 RUN tar xf /PaddleOCR/inference/{file}.tar -C /PaddleOCR/inference/
 
-# Download orc recognition model(light version). If you want to change normal version, you can change ch_ppocr_mobile_v1.1_rec_infer to ch_ppocr_server_v1.1_rec_infer, also remember change rec_model_dir in deploy/hubserving/ocr_system/params.py）
+# Download orc recognition model(light version). If you want to change normal version, you can change ch_ppocr_mobile_v2.0_rec_infer to ch_ppocr_server_v2.0_rec_infer, also remember change rec_model_dir in deploy/hubserving/ocr_system/params.py）
 ADD {link} /PaddleOCR/inference/
 RUN tar xf /PaddleOCR/inference/{file}.tar -C /PaddleOCR/inference/
 

diff --git a/deploy/docker/hubserving/gpu/Dockerfile b/deploy/docker/hubserving/gpu/Dockerfile
@@ -1,11 +1,9 @@
-# Version: 1.0.0
-FROM hub.baidubce.com/paddlepaddle/paddle:latest-gpu-cuda10.0-cudnn7-dev
+# Version: 2.0.0
+FROM egistry.baidubce.com/paddlepaddle/paddle:2.0.0rc1-gpu-cuda10.0-cudnn7
 
 # PaddleOCR base on Python3.7
 RUN pip3.7 install --upgrade pip -i https://mirror.baidu.com/pypi/simple
 
-RUN python3.7 -m pip install paddlepaddle-gpu==2.0.0rc0 -i https://mirror.baidu.com/pypi/simple
-
 RUN pip3.7 install paddlehub --upgrade -i https://mirror.baidu.com/pypi/simple
 
 RUN git clone https://github.com/PaddlePaddle/PaddleOCR.git /PaddleOCR
@@ -15,15 +13,15 @@ WORKDIR /PaddleOCR
 RUN pip3.7 install -r requirements.txt -i https://mirror.baidu.com/pypi/simple
 
 RUN mkdir -p /PaddleOCR/inference/
-# Download orc detect model(light version). if you want to change normal version, you can change ch_ppocr_mobile_v1.1_det_infer to ch_ppocr_server_v1.1_det_infer, also remember change det_model_dir in deploy/hubserving/ocr_system/params.py）
+# Download orc detect model(light version). if you want to change normal version, you can change ch_ppocr_mobile_v2.0_det_infer to ch_ppocr_server_v2.0_det_infer, also remember change det_model_dir in deploy/hubserving/ocr_system/params.py）
 ADD {link} /PaddleOCR/inference/
 RUN tar xf /PaddleOCR/inference/{file}.tar -C /PaddleOCR/inference/
 
-# Download direction classifier(light version). If you want to change normal version, you can change ch_ppocr_mobile_v1.1_cls_infer to ch_ppocr_mobile_v1.1_cls_infer, also remember change cls_model_dir in deploy/hubserving/ocr_system/params.py）
+# Download direction classifier(light version). If you want to change normal version, you can change ch_ppocr_mobile_v2.0_cls_infer to ch_ppocr_mobile_v2.0_cls_infer, also remember change cls_model_dir in deploy/hubserving/ocr_system/params.py）
 ADD {link} /PaddleOCR/inference/
 RUN tar xf /PaddleOCR/inference/{file} -C /PaddleOCR/inference/
 
-# Download orc recognition model(light version). If you want to change normal version, you can change ch_ppocr_mobile_v1.1_rec_infer to ch_ppocr_server_v1.1_rec_infer, also remember change rec_model_dir in deploy/hubserving/ocr_system/params.py）
+# Download orc recognition model(light version). If you want to change normal version, you can change ch_ppocr_mobile_v2.0_rec_infer to ch_ppocr_server_v2.0_rec_infer, also remember change rec_model_dir in deploy/hubserving/ocr_system/params.py）
 ADD {link} /PaddleOCR/inference/
 RUN tar xf /PaddleOCR/inference/{file}.tar -C /PaddleOCR/inference/