{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,9,17]],"date-time":"2025-09-17T06:10:43Z","timestamp":1758089443100,"version":"3.44.0"},"publisher-location":"Cham","reference-count":41,"publisher":"Springer Nature Switzerland","isbn-type":[{"type":"print","value":"9783032046239"},{"type":"electronic","value":"9783032046246"}],"license":[{"start":{"date-parts":[[2025,9,17]],"date-time":"2025-09-17T00:00:00Z","timestamp":1758067200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,9,17]],"date-time":"2025-09-17T00:00:00Z","timestamp":1758067200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2026]]},"DOI":"10.1007\/978-3-032-04624-6_30","type":"book-chapter","created":{"date-parts":[[2025,9,16]],"date-time":"2025-09-16T05:33:24Z","timestamp":1758000804000},"page":"509-525","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["VLMAWR: A Method for Manchu Archives Word Recognition Based on Vision-Language Model"],"prefix":"10.1007","author":[{"ORCID":"https:\/\/orcid.org\/0000-0001-5776-7083","authenticated-orcid":false,"given":"Yu","family":"Zhou","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0009-0003-8932-756X","authenticated-orcid":false,"given":"Zhengxu","family":"Jin","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-6942-8457","authenticated-orcid":false,"given":"Jianjun","family":"He","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0009-0002-0634-1183","authenticated-orcid":false,"given":"Baochun","family":"Wu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0009-0008-4557-8005","authenticated-orcid":false,"given":"Xinshu","family":"Cui","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-7490-6716","authenticated-orcid":false,"given":"Ruirui","family":"Zheng","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,9,17]]},"reference":[{"key":"30_CR1","unstructured":"Zhao, Y., Su, Y.: A review of the organization and research of Manchu archives since the 21st century, pp. 55\u201361, 73 (2017). (in Chinese)"},{"key":"30_CR2","unstructured":"Li, J.: Exploration of informatization of Manchu archives, pp. 44\u201345 (2019). (in Chinese)"},{"key":"30_CR3","unstructured":"Li, G.: From tradition to modernity: review and prospects of the compilation of Manchu archives at the First Historical Archives of China. (in Chinese)"},{"key":"30_CR4","doi-asserted-by":"publisher","unstructured":"Zhuang, X.: Research on the current status of Manchu document preservation in China. https:\/\/doi.org\/10.3969\/j.issn.1003-6547.2013.04.016. (in Chinese)","DOI":"10.3969\/j.issn.1003-6547.2013.04.016"},{"key":"30_CR5","doi-asserted-by":"publisher","unstructured":"Zhang, G., Li, J., Wang, A., Zhang, G., Li, J., Wang, A.: Extraction and recognition of offline handwritten Manchu stroke primitives. https:\/\/doi.org\/10.3969\/j.issn.1000-3428.2007.22.069. (in Chinese)","DOI":"10.3969\/j.issn.1000-3428.2007.22.069"},{"key":"30_CR6","doi-asserted-by":"publisher","unstructured":"Wei, W., Guo, C.: Off-line Manchu character recognition based on multi-classifier ensemble with combination features. https:\/\/doi.org\/10.3969\/j.issn.1000-7024.2012.06.050. (in Chinese)","DOI":"10.3969\/j.issn.1000-7024.2012.06.050"},{"key":"30_CR7","doi-asserted-by":"crossref","unstructured":"Shi B., Bai X., Yao C.: An end-to-end trainable neural network for image-based sequence recognition and its application to scene text recognition. IEEE Trans. Pattern Anal. Mach. Intel. 39(11), 2298\u20132304 (2016)","DOI":"10.1109\/TPAMI.2016.2646371"},{"key":"30_CR8","doi-asserted-by":"crossref","unstructured":"Graves, A., Fern\u00e1ndez, S., Gomez, F., et al.: Connectionist temporal classification: labelling un-segmented sequence data with recurrent neural networks. In: Proceedings of the 23rd International Conference on Machine Learning, pp. 369\u2013376 (2006)","DOI":"10.1145\/1143844.1143891"},{"key":"30_CR9","doi-asserted-by":"publisher","unstructured":"Li M., Zheng R., Xu S., Fu Y., Huang D.: Manchu Word Recognition Based on Convolutional Neural Network with Spatial Pyramid Pooling. https:\/\/doi.org\/10.1109\/CISP-BMEI.2018.8633131","DOI":"10.1109\/CISP-BMEI.2018.8633131"},{"key":"30_CR10","unstructured":"Gao, H.: Research on Manchu word recognition based on CADCN network and improved CTC loss. Master\u2019s Thesis, Dalian Minzu University (2023). (In Chinese)"},{"key":"30_CR11","doi-asserted-by":"crossref","unstructured":"Na B., Kim Y., Park S.: Multi-modal text recognition networks: Interactive enhancements between visual and semantic features. In: European conference on computer vision. Cham: Springer Nature Switzerland, pp.  446\u2013463 (2022)","DOI":"10.1007\/978-3-031-19815-1_26"},{"key":"30_CR12","doi-asserted-by":"crossref","unstructured":"Wang Y., Xie H., Fang S., Wang J., Zhu S., Zhang Y.: From two to one: a new scene text recognizer with visual language modeling network. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 14194\u201314203. (2021)","DOI":"10.1109\/ICCV48922.2021.01393"},{"key":"30_CR13","doi-asserted-by":"publisher","unstructured":"Fang S., Xie H., Wang Y., Mao Z., Zhang Y.: Read Like Humans: Autonomous, Bidirectional and Iterative Language Modeling for Scene Text Recognition. https:\/\/doi.org\/10.48550\/arXiv.2103.06495","DOI":"10.48550\/arXiv.2103.06495"},{"key":"30_CR14","unstructured":"Zhang, G., Li, J., He, R., Wang, A.: An offline recognition method of handwritten primitive Manchu characters based on strokes. In; Ninth International Workshop on Frontiers in Handwriting Recognition. IEEE, pp.  432\u2013437 (2004)"},{"key":"30_CR15","doi-asserted-by":"crossref","unstructured":"Zhang, G., Li, J., Wang, A.: A new recognition method for the handwritten Manchu character unit. In: 2006 International Conference on Machine Learning and Cybernetics. IEEE, pp. 3339\u20133344 (2006)","DOI":"10.1109\/ICMLC.2006.258471"},{"key":"30_CR16","first-page":"1061","volume":"11","author":"J Zhao","year":"2004","unstructured":"Zhao, J.: Manchu character recognition post-processing based on bayes rules and substitution set confusion matrix. J. Northeast. Univ. 11, 1061\u20131064 (2004). (in Chinese)","journal-title":"J. Northeast. Univ."},{"issue":"4","key":"30_CR17","first-page":"63","volume":"20","author":"J Zhao","year":"2006","unstructured":"Zhao, J., Li, J., Wang, L., et al.: Research on the post-processing of Manchu character recognition based on hidden Markov model. J. Chin. Inf. Process. 20(4), 63\u201367 (2006). (in Chinese)","journal-title":"J. Chin. Inf. Process."},{"issue":"1","key":"30_CR18","doi-asserted-by":"publisher","first-page":"43","DOI":"10.1080\/09720529.2016.1178912","volume":"20","author":"S Xu","year":"2017","unstructured":"Xu, S., Li, M., Zheng, R.R., et al.: Manchu character segmentation and recognition method. J. Discret. Math. Sci. Cryptogr. 20(1), 43\u201353 (2017). https:\/\/doi.org\/10.1080\/09720529.2016.1178912","journal-title":"J. Discret. Math. Sci. Cryptogr."},{"key":"30_CR19","doi-asserted-by":"publisher","unstructured":"Di, H., Min, L., Zheng, R., Shuang, X., Bi, J.: Synthetic data and DAG-SVM classifier for segmentation-free Manchu word recognition. In: Proceedings of the 2017 10th International Conference on Computational Intelligence and Security (CIIS), pp. 1\u20135 (2017). https:\/\/doi.org\/10.1109\/CIIS.2017.15","DOI":"10.1109\/CIIS.2017.15"},{"key":"30_CR20","doi-asserted-by":"publisher","unstructured":"Zheng, R., Xin, S., Zhou, Y., et al.: A K-shot Manchu recognition method for large category based on N-ary ECOC. J. Zhengzhou Univ. (Nat. Sci. Ed.) 53(4), 53\u201360 (2021). https:\/\/doi.org\/10.13705\/j.issn.1671-6841.2021178. (in Chinese)","DOI":"10.13705\/j.issn.1671-6841.2021178"},{"key":"30_CR21","doi-asserted-by":"publisher","unstructured":"Yang M., Wang D., Li Z., Yu X.: HMM inception-ResNet network classifier for Manchu word recognition. In: Proceedings of the 2022 2nd International Conference on Frontiers of Electronics, Information and Computation Technologies (ICFEICT), pp. 1\u20136. IEEE (2022). https:\/\/doi.org\/10.1109\/ICFEICT55092.2022.00010","DOI":"10.1109\/ICFEICT55092.2022.00010"},{"key":"30_CR22","doi-asserted-by":"crossref","unstructured":"Zheng, R., Li, M., He, J., Bi, J., ,Wu B.: Segmentation-free multi-font printed Manchu word recognition using deep convolutional features and data augmentation. In: 2018 11th International Congress on Image and Signal Processing, BioMedical Engineering and Informatics (CISP-BMEI). IEEE, pp. 1\u20136 (2018)","DOI":"10.1109\/CISP-BMEI.2018.8633208"},{"issue":"8","key":"30_CR23","first-page":"1","volume":"50","author":"S Xin","year":"2020","unstructured":"Xin, S., Zheng, R., Zhou, Y., et al.: A one-shot learning algorithm using support set information during training. J. Univ. Sci. Technol. China 50(8), 1 (2020). (in Chinese)","journal-title":"J. Univ. Sci. Technol. China"},{"key":"30_CR24","unstructured":"Fan, H.: Research on Manchu recognition method based on partial label learning and deep neural network. Master\u2019s Thesis, Dalian Minzu University (2022). (In Chinese)"},{"key":"30_CR25","doi-asserted-by":"publisher","unstructured":"Wang, Z., Lu, S., Wang, M., et al.: AMRE: an attention-based CRNN for Manchu word recognition on a woodblock-printed dataset. In: Proceedings of the International Conference on Neural Information Processing (ICONIP), pp. 267\u2013278. Springer, Cham (2022). https:\/\/doi.org\/10.1007\/978-3-031-30111-7_23","DOI":"10.1007\/978-3-031-30111-7_23"},{"key":"30_CR26","doi-asserted-by":"crossref","unstructured":"Li, H., Lu, S., Wei, X., Qi, Y.: Probing handwritten Manchu word recognition foundation model. In: Proceedings of the 2023 IEEE 29th International Conference on Parallel and Distributed Systems (ICPADS), pp. 1262\u20131269 (2023)","DOI":"10.1109\/ICPADS60453.2023.00182"},{"key":"30_CR27","doi-asserted-by":"publisher","unstructured":"Tan, Z., Li, M., Wang, D.: Research on Manchu recognition based on multi-scale feature fusion swin transformer. J. Jilin Normal Univ. (Nat. Sci. Ed.) 46(1), 103\u2013110 (2025). https:\/\/doi.org\/10.16862\/j.cnki.issn1674-3873.2025.01.015. (in Chinese)","DOI":"10.16862\/j.cnki.issn1674-3873.2025.01.015"},{"key":"30_CR28","doi-asserted-by":"publisher","unstructured":"Liu, H., Jin, S., Zhang, C.: Connectionist temporal classification with maximum entropy regularization. In: Advances in Neural Information Processing Systems, vol. 31, pp. 1\u201310 (2018). https:\/\/doi.org\/10.48550\/arXiv.1805.04908","DOI":"10.48550\/arXiv.1805.04908"},{"key":"30_CR29","doi-asserted-by":"crossref","unstructured":"Xie, Z., Huang, Y., Zhu, Y., et al.: Aggregation cross-entropy for sequence recognition. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. Long Beach, USA, pp. 6538\u20136547 (2019)","DOI":"10.1109\/CVPR.2019.00670"},{"key":"30_CR30","doi-asserted-by":"crossref","unstructured":"Hu, W., Cai, X., Hou, J., et al.: GTC: guided training of CTC towards efficient and accurate scene text recognition. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 34, no. 07, pp. 11005\u201311012 (2020)","DOI":"10.1609\/aaai.v34i07.6735"},{"key":"30_CR31","doi-asserted-by":"crossref","unstructured":"Du, Y., Chen, Z., Jia, C., et al.: SVTR: Scene text recognition with a single visual model. arXiv preprint arXiv:2205.00159 (2022)","DOI":"10.24963\/ijcai.2022\/124"},{"key":"30_CR32","unstructured":"Jaderberg, M., Simonyan, K., Vedaldi, A., et al.: Deep structured out-put learning for unconstrained text recognition. arXiv preprint arXiv:1412.5903 (2014)"},{"key":"30_CR33","doi-asserted-by":"crossref","unstructured":"Lee, C.Y., Osindero, S.: Recursive recurrent nets with attention modeling for OCR in the wild. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (CVPR), pp. 2231\u20132239. IEEE Press (2016)","DOI":"10.1109\/CVPR.2016.245"},{"key":"30_CR34","doi-asserted-by":"crossref","unstructured":"Qiao, Z., Zhou, Y., Yang, D., et al.: SEED: semantics enhanced encoder-decoder framework for scene text recognition. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp. 13528\u201313537. IEEE Press (2020)","DOI":"10.1109\/CVPR42600.2020.01354"},{"key":"30_CR35","doi-asserted-by":"crossref","unstructured":"Shi, B., Wang, X., Lyu, P., et al.: Robust scene text recognition with automatic rectification. In: Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (CVPR), pp. 4168\u20134176. IEEE Press (2016)","DOI":"10.1109\/CVPR.2016.452"},{"issue":"9","key":"30_CR36","doi-asserted-by":"publisher","first-page":"2035","DOI":"10.1109\/TPAMI.2018.2848939","volume":"41","author":"B Shi","year":"2018","unstructured":"Shi, B., Yang, M., Wang, X., et al.: ASTER: an attentional scene text recognizer with flexible rectification. IEEE Trans. Pattern Anal. Mach. Intell. 41(9), 2035\u20132048 (2018)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"30_CR37","doi-asserted-by":"crossref","unstructured":"Yu, D., Li, X., Zhang, C., et al.: Towards accurate scene text recognition with semantic reasoning networks. In: Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp. 12113\u201312122. IEEE Press (2020)","DOI":"10.1109\/CVPR42600.2020.01213"},{"key":"30_CR38","doi-asserted-by":"crossref","unstructured":"Wojna, Z., Gorban, A.N., Lee, D.S., et al.: Attention-based extraction of structured information from street view imagery. In: Proceedings of 14th International Conference on Document Analysis and Recognition (ICDAR), vol. 1, pp. 844\u2013850. IEEE, Cham (2017)","DOI":"10.1109\/ICDAR.2017.143"},{"key":"30_CR39","doi-asserted-by":"crossref","unstructured":"Wang, T., Zhu, Y., Jin, L., et al.: Decoupled attention network for text recognition. In: Proceedings of AAAI Conference on Artificial Intelligence, vol. 34, no. 7, pp. 12216\u201312224. AAAI Press (2020)","DOI":"10.1609\/aaai.v34i07.6903"},{"key":"30_CR40","doi-asserted-by":"publisher","unstructured":"Tsai, Y.H.H., Bai, S., Liang, P.P., et al.: Multimodal transformer for unaligned multimodal language sequences. In: Proceedings of the 57th Annual Meeting of the Association for Computational Linguistics (ACL), pp. 6558\u20136569 (2019). https:\/\/doi.org\/10.18653\/v1\/P19-1656","DOI":"10.18653\/v1\/P19-1656"},{"key":"30_CR41","doi-asserted-by":"publisher","unstructured":"Wang, P., Da, C., Yao, C.: Multi-granularity prediction for scene text recognition. In: Computer Vision\u2013ECCV 2022: 17th European Conference, Tel Aviv, Israel, 23\u201327 October 2022, Proceedings, Part XXVIII, pp. 339\u2013355. Springer, Cham (2022). https:\/\/doi.org\/10.1007\/978-3-031-25066-8_20","DOI":"10.1007\/978-3-031-25066-8_20"}],"container-title":["Lecture Notes in Computer Science","Document Analysis and Recognition \u2013 ICDAR 2025"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-032-04624-6_30","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,16]],"date-time":"2025-09-16T05:33:57Z","timestamp":1758000837000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-032-04624-6_30"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,9,17]]},"ISBN":["9783032046239","9783032046246"],"references-count":41,"URL":"https:\/\/doi.org\/10.1007\/978-3-032-04624-6_30","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"type":"print","value":"0302-9743"},{"type":"electronic","value":"1611-3349"}],"subject":[],"published":{"date-parts":[[2025,9,17]]},"assertion":[{"value":"17 September 2025","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Disclosure of Interests"}},{"value":"ICDAR","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"International Conference on Document Analysis and Recognition","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Wuhan","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"China","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2025","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"16 September 2025","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"21 September 2025","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"18","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"icdar2025","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/iapr.org\/icdar2025","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}