{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,1,24]],"date-time":"2026-01-24T16:39:17Z","timestamp":1769272757983,"version":"3.49.0"},"reference-count":49,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2025,4,16]],"date-time":"2025-04-16T00:00:00Z","timestamp":1744761600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,4,16]],"date-time":"2025-04-16T00:00:00Z","timestamp":1744761600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Jiangsu Provincial Education Planning Project","award":["B-b\/2024\/01\/01"],"award-info":[{"award-number":["B-b\/2024\/01\/01"]}]},{"name":"Special Research Project on the Practice of Digital Transformation and Modernization in Higher Education in Jiangsu Province","award":["2024CXJG100"],"award-info":[{"award-number":["2024CXJG100"]}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62101245"],"award-info":[{"award-number":["62101245"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["SIViP"],"published-print":{"date-parts":[[2025,6]]},"DOI":"10.1007\/s11760-025-04058-y","type":"journal-article","created":{"date-parts":[[2025,4,16]],"date-time":"2025-04-16T13:56:56Z","timestamp":1744811816000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Enhanced Chinese scene text recognition model base on cross-domain feature fusion"],"prefix":"10.1007","volume":"19","author":[{"given":"Ran","family":"Cui","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Aichun","family":"Zhu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zichen","family":"Ding","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,4,16]]},"reference":[{"key":"4058_CR1","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1016\/j.engappai.2023.106271","volume":"123","author":"Y Guo","year":"2023","unstructured":"Guo, Y., Yu, H., Ma, L., Zeng, L., Luo, X.: THFE: a triple-hierarchy feature enhancement method for tiny boat detection. Eng. Appl. Artif. Intell. 123, 1\u201316 (2023)","journal-title":"Eng. Appl. Artif. Intell."},{"key":"4058_CR2","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1016\/j.patcog.2024.110329","volume":"150","author":"Y Guo","year":"2024","unstructured":"Guo, Y., Yu, H., Xie, S., Ma, L., Cao, X., Luo, X.: DSCA: a dual semantic correlation alignment method for domain adaptation object detection. Pattern Recognit. 150, 1\u201313 (2024)","journal-title":"Pattern Recognit."},{"key":"4058_CR3","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1016\/j.knosys.2024.111964","volume":"297","author":"Y Guo","year":"2024","unstructured":"Guo, Y., Ma, L., Luo, X., Xie, S.: DP-DDCL: a discriminative prototype with dual decoupled contrast learning method for few-shot object detection. Knowl. Based Syst. 297, 1\u201315 (2024)","journal-title":"Knowl. Based Syst."},{"issue":"11","key":"4058_CR4","doi-asserted-by":"publisher","first-page":"10646","DOI":"10.1109\/TCSVT.2024.3407057","volume":"34","author":"Y Guo","year":"2024","unstructured":"Guo, Y., Yu, H., Ma, L., Luo, X., Xie, S.: DIE-CDK: a discriminative information enhancement method with cross-modal domain knowledge for fine-grained ship detection. IEEE Trans. Circuits Syst. Video Technol. 34(11), 10646\u201310661 (2024)","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"4058_CR5","doi-asserted-by":"crossref","unstructured":"Sermanet, P., Lecun, Y.: Traffic sign recognition with multi-scale convolutional networks. In: Proceedings of the International Joint Conference on Neural Networks, pp. 2809\u20132813 (2011)","DOI":"10.1109\/IJCNN.2011.6033589"},{"key":"4058_CR6","doi-asserted-by":"crossref","unstructured":"Wang, Z., Liu, X., Li, H., Sheng, L., Yan, J., Wang, X., Shao, J.: Camp: Cross-modal adaptive message passing for text-image retrieval. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 5763\u20135772 (2019)","DOI":"10.1109\/ICCV.2019.00586"},{"key":"4058_CR7","unstructured":"Yu, H., Chen, J., Li, B., Ma, J., Guan, M., Xu, X., Wang, X., Qu, S., Xue, X.: Benchmarking Chinese text recognition: datasets, baselines, and an empirical study. arXiv (2021)"},{"key":"4058_CR8","unstructured":"Yin, F., Wu, Y.-C., Zhang, X.-Y., Liu, C.-L.: Scene text recognition with sliding convolutional character models. arXiv preprint arXiv:1709.01727 (2017)"},{"key":"4058_CR9","doi-asserted-by":"crossref","unstructured":"Liu, W., Chen, C., Wong, K.-Y.: Char-net: a character-aware neural network for distorted scene text recognition. In: Proceedings of the AAAI Conference on Artificial Intelligence (2018)","DOI":"10.1609\/aaai.v32i1.12246"},{"issue":"11","key":"4058_CR10","doi-asserted-by":"publisher","first-page":"2298","DOI":"10.1109\/TPAMI.2016.2646371","volume":"39","author":"B Shi","year":"2016","unstructured":"Shi, B., Bai, X., Yao, C.: An end-to-end trainable neural network for image-based sequence recognition and its application to scene text recognition. IEEE Trans. Pattern Anal. Mach. Intell. 39(11), 2298\u20132304 (2016)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"9","key":"4058_CR11","doi-asserted-by":"publisher","first-page":"2035","DOI":"10.1109\/TPAMI.2018.2848939","volume":"41","author":"B Shi","year":"2018","unstructured":"Shi, B., Yang, M., Wang, X., Lyu, P., Yao, C., Bai, X.: ASTER: an attentional scene text recognizer with flexible rectification. IEEE Trans. Pattern Anal. Mach. Intell. 41(9), 2035\u20132048 (2018)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"4058_CR12","doi-asserted-by":"crossref","unstructured":"Li, H., Wang, P., Shen, C., Zhang, G.: Show, attend and read: A simple and strong baseline for irregular text recognition. In: Proceedings of the AAAI Conference on Artificial Intelligence, pp. 8610\u20138617 (2019)","DOI":"10.1609\/aaai.v33i01.33018610"},{"key":"4058_CR13","first-page":"1","volume":"63","author":"Z Li","year":"2025","unstructured":"Li, Z., Hu, J., Wu, K., Miao, J., Wu, J.: Comprehensive attribute difference attention network for remote sensing image semantic understanding. IEEE Trans. Geosci. Remote Sens. 63, 1\u201316 (2025)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4058_CR14","doi-asserted-by":"publisher","first-page":"12597","DOI":"10.1038\/s41598-024-63363-7","volume":"14","author":"Z Li","year":"2024","unstructured":"Li, Z., Hu, J., Wu, K., et al.: Local feature acquisition and global context understanding network for very high-resolution land cover classification. Sci. Rep. 14, 12597 (2024)","journal-title":"Sci. Rep."},{"key":"4058_CR15","first-page":"1","volume":"62","author":"Z Li","year":"2024","unstructured":"Li, Z., Hu, J., Wu, K., Miao, J., Wu, J.: Adjacent-Atrous mechanism for expanding global receptive fields: An end-to-end network for multiattribute scene analysis in remote sensing imagery. IEEE Trans. Geosci. Remote Sens. 62, 1\u201319 (2024)","journal-title":"IEEE Trans. Geosci. Remote Sens."},{"key":"4058_CR16","doi-asserted-by":"crossref","unstructured":"Fang, S., Xie, H., Wang, Y., Mao, Z., Zhang, Y.: Read like humans: Autonomous, bidirectional and iterative language modeling for scene text recognition. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7098\u20137107 (2021)","DOI":"10.1109\/CVPR46437.2021.00702"},{"key":"4058_CR17","unstructured":"Wang, K., Babenko, B., Belongie, S.: End-to-end scene text recognition. In: Proceedings of the IEEE International Conference on Computer Vision (ICCV), pp. 1457\u20131464 (2011)"},{"key":"4058_CR18","unstructured":"Wang, T., Wu, D.J., Coates, A., Ng, A.Y.: End-to-end text recognition with convolutional neural networks. In: Proceedings of the 21st International Conference on Pattern Recognition (ICPR), pp. 3304\u20133308 (2012)"},{"key":"4058_CR19","doi-asserted-by":"crossref","unstructured":"Liu, X., Liang, D., Yan, S., Chen, D., Qiao, Y., Mei, T.: Scene text recognition with convolutional neural network and weight finite state transducer. In: Proceedings of the 23rd International Conference on Pattern Recognition (ICPR), pp. 3999\u20134004 (2016)","DOI":"10.1109\/ICPR.2016.7900259"},{"key":"4058_CR20","doi-asserted-by":"publisher","first-page":"30","DOI":"10.1016\/j.cviu.2016.01.002","volume":"154","author":"A Mishra","year":"2016","unstructured":"Mishra, A., Alahari, K., Jawahar, C.V.: Enhancing energy minimization framework for scene text recognition with top-down cues. Comput. Vis. Image Underst. 154, 30\u201342 (2016)","journal-title":"Comput. Vis. Image Underst."},{"key":"4058_CR21","doi-asserted-by":"crossref","unstructured":"Phan, T.Q., Shivakumara, P., Tian, S., Lim, S.K., Tan, C.L.: Recognizing text with perspective distortion in natural scenes. In: Proceedings of the IEEE International Conference on Computer Vision (ICCV), pp. 569\u2013576 (2013)","DOI":"10.1109\/ICCV.2013.76"},{"key":"4058_CR22","doi-asserted-by":"crossref","unstructured":"Yao, C., Bai, X., Liu, W.: Strokelets: a learned multi-scale representation for scene text recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR), pp. 4042\u20134049 (2014)","DOI":"10.1109\/CVPR.2014.515"},{"key":"4058_CR23","doi-asserted-by":"publisher","first-page":"45","DOI":"10.1007\/s10032-023-00444-9","volume":"27","author":"S Wu","year":"2024","unstructured":"Wu, S., Li, Y., Wang, Z.: Chinese text recognition enhanced by glyph and character semantic information. IJDAR 27, 45\u201356 (2024)","journal-title":"IJDAR"},{"key":"4058_CR24","doi-asserted-by":"crossref","unstructured":"Gordo, A.: Supervised mid-level features for word image representation, pp. 2956\u20132964 (2015)","DOI":"10.1109\/CVPR.2015.7298914"},{"key":"4058_CR25","doi-asserted-by":"crossref","unstructured":"Baek, J., Kim, G., Lee, J., Lee, S.: What is wrong with scene text recognition model comparisons? dataset and model analysis. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 4704\u20134713 (2019)","DOI":"10.1109\/ICCV.2019.00481"},{"issue":"11","key":"4058_CR26","doi-asserted-by":"publisher","first-page":"2298","DOI":"10.1109\/TPAMI.2016.2646371","volume":"39","author":"B Shi","year":"2016","unstructured":"Shi, B., Bai, X., Yao, C.: An end-to-end trainable neural network for image-based sequence recognition and its application to scene text recognition. IEEE Trans. Pattern Anal. Mach. Intell. 39(11), 2298\u20132304 (2016)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"4058_CR27","unstructured":"Simonyan, K., Zisserman, A.: Very deep convolutional networks for large-scale image recognition (2015)"},{"key":"4058_CR28","doi-asserted-by":"crossref","unstructured":"Cheng, Z., Bai, F., Xu, Y., Zheng, G., Pu, S., Zhou, S.: Focusing attention: towards accurate text recognition in natural images. In: Proceedings of the IEEE International Conference on Computer Vision (ICCV) (2017)","DOI":"10.1109\/ICCV.2017.543"},{"key":"4058_CR29","doi-asserted-by":"crossref","unstructured":"Zhu, Y., Wang, S., Huang, Z., Chen, K.: Text recognition in images based on transformer with hierarchical attention. In: 2019 IEEE International Conference on Image Processing, pp. 1945\u20131949 (2019)","DOI":"10.1109\/ICIP.2019.8803203"},{"key":"4058_CR30","unstructured":"Lyu, P., Zhang, C., Liu, S., Qiao, M., Xu, Y., Wu, L., Yao, K., Han, J., Ding, E., Wang, J.: Maskocr: Text recognition with masked encoder-decoder pretraining. arXiv (2023)"},{"issue":"4","key":"4058_CR31","doi-asserted-by":"publisher","first-page":"960","DOI":"10.1007\/s11263-020-01411-1","volume":"129","author":"C Luo","year":"2021","unstructured":"Luo, C., Jin, L., Sun, Z.: Separating content from style using adversarial learning for recognizing text in the wild. Int. J. Comput. Vis. 129(4), 960\u2013976 (2021)","journal-title":"Int. J. Comput. Vis."},{"key":"4058_CR32","unstructured":"Goodfellow, I., Pouget-Abadie, J., Mirza, M., Xu, B., Warde-Farley, D., Ozair, S., Bengio, Y.: Generative adversarial nets. Adv. Neural Inf. Process. Syst. 27(4) (2014)"},{"key":"4058_CR33","doi-asserted-by":"crossref","unstructured":"Wang, W., Xie, E., Li, X., Hou, W., Lu, T., Yu, G., Zhang, Y.: Scene text image super-resolution in the wild. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 1\u201317 (2020)","DOI":"10.1007\/978-3-030-58607-2_38"},{"key":"4058_CR34","doi-asserted-by":"crossref","unstructured":"Mou, Y., Wang, Z., Xu, H., Lu, T., Xiang, Y.: Plugnet: Degradation aware scene text recognition supervised by a pluggable super-resolution unit. In: Proceedings of the European Conference on Computer Vision (ECCV) (2020)","DOI":"10.1007\/978-3-030-58555-6_10"},{"key":"4058_CR35","doi-asserted-by":"crossref","unstructured":"Zhan, F., Zhu, H., Lu, S.: Spatial fusion gan for image synthesis. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2019)","DOI":"10.1109\/CVPR.2019.00377"},{"key":"4058_CR36","doi-asserted-by":"crossref","unstructured":"Yang, M., Luo, H., Xu, Z., Chen, C., Wang, J.: Symmetry-constrained rectification network for scene text recognition. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV) (2019)","DOI":"10.1109\/ICCV.2019.00924"},{"key":"4058_CR37","doi-asserted-by":"crossref","unstructured":"Sheng, F., Chen, Z., Xu, B.: NRTR: a no-recurrence sequence-to-sequence model for scene text recognition. In: 2019 International Conference on Document Analysis and Recognition, pp. 781\u2013786 (2019)","DOI":"10.1109\/ICDAR.2019.00130"},{"issue":"14","key":"4058_CR38","doi-asserted-by":"publisher","first-page":"261","DOI":"10.1016\/j.neucom.2019.11.049","volume":"381","author":"X Chen","year":"2020","unstructured":"Chen, X., Wang, T., Zhu, Y., Jin, L., Luo, C.: Adaptive embedding gate for attention-based scene text recognition. Neurocomputing 381(14), 261\u2013271 (2020)","journal-title":"Neurocomputing"},{"key":"4058_CR39","doi-asserted-by":"crossref","unstructured":"Yu, D., Li, X., Zhang, C., Liu, T., Han, J., Liu, J., Ding, E.: Towards accurate scene text recognition with semantic reasoning networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12113\u201312122 (2020)","DOI":"10.1109\/CVPR42600.2020.01213"},{"key":"4058_CR40","doi-asserted-by":"crossref","unstructured":"Fang, S., Xie, H., Wang, Y., Mao, Z., Zhang, Y.: Read like humans: autonomous, bidirectional and iterative language modeling for scene text recognition. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition(CVPR), pp. 7098\u20137107 (2021)","DOI":"10.1109\/CVPR46437.2021.00702"},{"key":"4058_CR41","doi-asserted-by":"crossref","unstructured":"Du, Y., Chen, Z., Jia, C., Yin, X., Zheng, T., Li, C., Du, Y., Jiang, Y.-G.: SVTR: scene text recognition with a single visual model. arXiv (2022)","DOI":"10.24963\/ijcai.2022\/124"},{"key":"4058_CR42","unstructured":"Chen, J., Yu, H., Ma, J., Guan, M., Xu, X., Wang, X., Qu, S., Li, B., Xue, X.: Benchmarking Chinese text recognition: datasets, baselines, and an empirical study. arXiv (2021)"},{"key":"4058_CR43","doi-asserted-by":"crossref","unstructured":"Shi, B., Yao, C., Liao, M., Yang, M., Xu, P., Cui, L., Belongie, S., Lu, S., Bai, X.: ICDAR2017 competition on reading Chinese text in the wild (RCTW-17). In: ICDAR (2017)","DOI":"10.1109\/ICDAR.2017.233"},{"key":"4058_CR44","doi-asserted-by":"crossref","unstructured":"Zhang, R., Zhou, Y., Jiang, Q., Song, Q., Li, N., Zhou, K., Wang, L., Wang, D., Liao, M., Yang, M., et al.: ICDAR 2019 robust reading challenge on reading Chinese text on signboard. In: ICDAR (2019)","DOI":"10.1109\/ICDAR.2019.00253"},{"key":"4058_CR45","doi-asserted-by":"crossref","unstructured":"Sun, Y., Ni, Z., Chng, C.-K., Liu, Y., Luo, C., Ng, C.C., Han, J., Ding, E., Liu, J., Karatzas, D., et al.: ICDAR 2019 competition on large-scale street view text with partial labeling-RRC-LSVT. In: ICDAR (2019)","DOI":"10.1109\/ICDAR.2019.00250"},{"key":"4058_CR46","doi-asserted-by":"crossref","unstructured":"Chng, C.K., Liu, Y., Sun, Y., Ng, C.C., Luo, C., Ni, Z., Fang, C., Zhang, S., Han, J., Ding, E., et al.: ICDAR2019 robust reading challenge on arbitrary-shaped text-RRC-ArT. In: ICDAR (2019)","DOI":"10.1109\/ICDAR.2019.00252"},{"key":"4058_CR47","doi-asserted-by":"publisher","first-page":"509","DOI":"10.1007\/s11390-019-1923-y","volume":"34","author":"T-L Yuan","year":"2019","unstructured":"Yuan, T.-L., Zhu, Z., Xu, K., Li, C.-J., Mu, T.-J., Hu, S.-M.: A large Chinese text dataset in the wild. J. Comput. Sci. Technol. 34, 509 (2019)","journal-title":"J. Comput. Sci. Technol."},{"key":"4058_CR48","doi-asserted-by":"crossref","unstructured":"Li, H., Wang, P., Shen, C., Zhang, G.: Show, attend and read: A simple and strong baseline for irregular text recognition. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 33, pp. 8610\u20138617 (2019)","DOI":"10.1609\/aaai.v33i01.33018610"},{"key":"4058_CR49","doi-asserted-by":"crossref","unstructured":"Qiao, Z., Zhou, Y., Yang, D., Zhou, Y., Wang, W.: Seed: Semantics enhanced encoder-decoder framework for scene text recognition. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition(CVPR), pp. 13528\u201313537 (2020)","DOI":"10.1109\/CVPR42600.2020.01354"}],"container-title":["Signal, Image and Video Processing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11760-025-04058-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11760-025-04058-y\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11760-025-04058-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,5,19]],"date-time":"2025-05-19T10:38:50Z","timestamp":1747651130000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11760-025-04058-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,4,16]]},"references-count":49,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2025,6]]}},"alternative-id":["4058"],"URL":"https:\/\/doi.org\/10.1007\/s11760-025-04058-y","relation":{},"ISSN":["1863-1703","1863-1711"],"issn-type":[{"value":"1863-1703","type":"print"},{"value":"1863-1711","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,4,16]]},"assertion":[{"value":"20 November 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 March 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"11 March 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"16 April 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no Conflict of interest. Specifically, none of the authors have any financial or personal relationships with other people or organizations that could inappropriately influence (bias) their work. In addition, no authors have received any funding or grants from any agency or organization that could give rise to potential Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"499"}}