{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,9,26]],"date-time":"2025-09-26T15:40:10Z","timestamp":1758901210845,"version":"3.44.0"},"reference-count":46,"publisher":"Springer Science and Business Media LLC","issue":"32","license":[{"start":{"date-parts":[[2025,3,15]],"date-time":"2025-03-15T00:00:00Z","timestamp":1741996800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0"},{"start":{"date-parts":[[2025,3,15]],"date-time":"2025-03-15T00:00:00Z","timestamp":1741996800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0"}],"funder":[{"DOI":"10.13039\/100020595","name":"National Science and Technology Council","doi-asserted-by":"publisher","award":["110-2221-E-004-008-MY3"],"award-info":[{"award-number":["110-2221-E-004-008-MY3"]}],"id":[{"id":"10.13039\/100020595","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-025-20736-y","type":"journal-article","created":{"date-parts":[[2025,3,15]],"date-time":"2025-03-15T02:11:48Z","timestamp":1742004708000},"page":"39375-39397","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Detecting reading orders with block-based models for OCR for classical Chinese documents: concepts and demonstrations"],"prefix":"10.1007","volume":"84","author":[{"given":"Hsing-Yuan","family":"Ma","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-4093-1497","authenticated-orcid":false,"given":"Chao-Lin","family":"Liu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hen-Hsen","family":"Huang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,3,15]]},"reference":[{"doi-asserted-by":"publisher","unstructured":"Ceci M, Berardi M, Porcelli G, Malerba D (2007) A data mining approach to reading order detection. In: Proceedings of the 9th international conference on document analysis and recognition, vol 2, pp 924\u2013928. https:\/\/doi.org\/10.1109\/ICDAR.2007.4377050","key":"20736_CR1","DOI":"10.1109\/ICDAR.2007.4377050"},{"doi-asserted-by":"publisher","unstructured":"Malerba D, Ceci M, Berardi M (2007) Machine learning for reading order detection in document image understanding. In: Machine learning in document analysis and recognition. studies in computational intelligence, vol 90, pp 45\u201369. https:\/\/doi.org\/10.1007\/978-3-540-76280-5_3","key":"20736_CR2","DOI":"10.1007\/978-3-540-76280-5_3"},{"doi-asserted-by":"publisher","unstructured":"Malerba D, Ceci M (2008) Learning to order: A relational approach. In: Ra\u015b, ZW, Tsumoto S, Zighed D (eds) Mining complex data, pp 209\u2013223. Springer, Berlin, Heidelberg. https:\/\/doi.org\/10.1007\/978-3-540-68416-9_17","key":"20736_CR3","DOI":"10.1007\/978-3-540-68416-9_17"},{"doi-asserted-by":"publisher","unstructured":"Ferilli S, Grieco D, Redavid D, Esposito F (2014) Abstract argumentation for reading order detection. In: Proceedings of the ACM symposium on document engineering, pp 45\u201348.https:\/\/doi.org\/10.1145\/2644866.2644883","key":"20736_CR4","DOI":"10.1145\/2644866.2644883"},{"doi-asserted-by":"publisher","unstructured":"Li L, Gao F, Bu J, Wang Y, Yu Z, Zheng Q (2020) An end-to-end OCR text re-organization sequence learning for rich-text detail image comprehension. In: Proceedings of the 16th european conference on computer vision, pp 85\u2013100. https:\/\/doi.org\/10.1007\/978-3-030-58595-2_6","key":"20736_CR5","DOI":"10.1007\/978-3-030-58595-2_6"},{"key":"20736_CR6","doi-asserted-by":"publisher","first-page":"9593","DOI":"10.1007\/s00521-022-06948-5","volume":"34","author":"L Quir\u00f3s","year":"2022","unstructured":"Quir\u00f3s L, Vidal E (2022) Reading order detection on handwritten documents. Neural Comput Appl 34:9593\u20139611. https:\/\/doi.org\/10.1007\/s00521-022-06948-5","journal-title":"Neural Comput Appl"},{"issue":"9","key":"20736_CR7","doi-asserted-by":"publisher","first-page":"767","DOI":"10.1080\/08839510600903858","volume":"20","author":"M Aiello","year":"2006","unstructured":"Aiello M, Pegoretti A (2006) Textual article clustering in newspaper pages. Appl Artif Intell 20(9):767\u2013796. https:\/\/doi.org\/10.1080\/08839510600903858","journal-title":"Appl Artif Intell"},{"issue":"2","key":"20736_CR8","doi-asserted-by":"publisher","first-page":"485","DOI":"10.1016\/S0031-3203(01)00026-7","volume":"35","author":"J-Y Lee","year":"2002","unstructured":"Lee J-Y, Park J-S, Byun H, Moon J, Lee S-W (2002) Automatic generation of structured hyperdocuments from document images. Pattern Recognit 35(2):485\u2013503. https:\/\/doi.org\/10.1016\/S0031-3203(01)00026-7","journal-title":"Pattern Recognit"},{"unstructured":"Breuel TM (2003) High performance document layout analysis. In: Proceedings of the 2003 symposium on document image understanding technology, pp 209\u2013218","key":"20736_CR9"},{"doi-asserted-by":"publisher","unstructured":"Quir\u00f3s L, Vidal E (2020) Learning to sort handwritten text lines in reading order through estimated binary order relations. In: Proceedings of the 25th international conference on pattern recognition, pp 7661\u20137668. https:\/\/doi.org\/10.1109\/ICPR48806.2021.9413256","key":"20736_CR10","DOI":"10.1109\/ICPR48806.2021.9413256"},{"doi-asserted-by":"publisher","unstructured":"Wang Z, Xu Y, Cui L, Shang J, Wei F (2021) LayoutReader: Pre-training of text and layout for reading order detection. In: Proceedings of the 2021 conference on empirical methods in natural language processing, pp 4735\u20134744. https:\/\/doi.org\/10.18653\/v1\/2021.emnlp-main.389","key":"20736_CR11","DOI":"10.18653\/v1\/2021.emnlp-main.389"},{"doi-asserted-by":"publisher","unstructured":"Clausner C, Pletschacher S, Antonacopoulos A (2013) The significance of reading order in document recognition and its evaluation. In: Proceedings of the 12th international conference on document analysis and recognition, pp 688\u2013692. https:\/\/doi.org\/10.1109\/ICDAR.2013.141","key":"20736_CR12","DOI":"10.1109\/ICDAR.2013.141"},{"doi-asserted-by":"publisher","unstructured":"Tang C-W, Liu C-L, Chiu P-S (2021) HRRegionNet: Chinese character segmentation in historical documents with regional awareness. In: Lecture notes in computer science 12824: proceedings of the 16th international conference on document analysis and recognition, vol 4, pp 3\u201317. https:\/\/doi.org\/10.1007\/978-3-030-86337-1_1","key":"20736_CR13","DOI":"10.1007\/978-3-030-86337-1_1"},{"doi-asserted-by":"publisher","unstructured":"Tang C-W, Liu C-L, Chiu P-S (2020) HRCenterNet: An anchorless approach to Chinese character segmentation in historical documents. In: Proceedings of the 5th workshop on computational archival science: digital records in the age of big data, 2020 IEEE international conference on big data, pp 1924\u20131930.https:\/\/doi.org\/10.1109\/BigData50022.2020.9378051","key":"20736_CR14","DOI":"10.1109\/BigData50022.2020.9378051"},{"doi-asserted-by":"publisher","unstructured":"Zhuang L, Bao T, Zhu X, Wang C, Naoi S (2004) A Chinese OCR spelling check approach based on statistical language models. In: Proceedings of the 2004 IEEE international conference on systems, man and cybernetics, pp 4727\u20134732. https:\/\/doi.org\/10.1109\/ICSMC.2004.1401278","key":"20736_CR15","DOI":"10.1109\/ICSMC.2004.1401278"},{"doi-asserted-by":"publisher","unstructured":"Wang H-A, Liu P-T (2019) Towards a higher accuracy of optical character recognition of Chinese rare books in making use of text model. In: Proceedings of the 3rd international conference on digital access to textual cultural heritage, pp 15\u201318. https:\/\/doi.org\/10.1145\/3322905.3322922","key":"20736_CR16","DOI":"10.1145\/3322905.3322922"},{"unstructured":"Simske S, Vans M (2021) Functional Applications of Text Analytics Systems. Routledge, London and New York. https:\/\/www.routledge.com\/Functional-Applications-of-Text-Analytics-Systems\/Simske-Vans\/p\/book\/9788770223430","key":"20736_CR17"},{"unstructured":"Luo W (2004) Concise Organization and Version Study of Ancient Books. Macao Library & Information Management Association, China:Macao. http:\/\/mlima.org.mo\/www\/mlima\/mlimaInfoShow_6266","key":"20736_CR18"},{"unstructured":"Liu, Z-Y (2007) Understanding of Printed Ancient Book and Book Collectors. Taiwan Student Bookstore, Taiwan:Taipei. http:\/\/www.studentbook.com.tw\/","key":"20736_CR19"},{"key":"20736_CR20","doi-asserted-by":"publisher","first-page":"17209","DOI":"10.1007\/s00521-020-05563-6","volume":"32","author":"J Mart\u00ednek","year":"2020","unstructured":"Mart\u00ednek J, Lenc L, Kr\u00e1l P (2020) Building an efficient ocr system for historical documents with little training data. Neural Comput Appl 32:17209\u201317227. https:\/\/doi.org\/10.1007\/s00521-020-05563-6","journal-title":"Neural Comput Appl"},{"doi-asserted-by":"publisher","unstructured":"Ma H-Y, Huang H-H, Liu C-L (2024) Reading between the lines: Image-based order detection in OCR for Chinese historical documents. In: Proceedings of the 38th annual AAAI conference on artificial intelligence, pp 23808\u201323810. https:\/\/doi.org\/10.1145\/3394486.3403172","key":"20736_CR21","DOI":"10.1145\/3394486.3403172"},{"issue":"1","key":"20736_CR22","doi-asserted-by":"publisher","first-page":"19","DOI":"10.6129\/CJP.20160304","volume":"58","author":"C-H Chen","year":"2016","unstructured":"Chen C-H, Tsai J-L (2016) Eye movement evidence for the effects of word boundary cue when reading Chinese. Chin J Psychol 58(1):19\u201344. https:\/\/doi.org\/10.6129\/CJP.20160304","journal-title":"Chin J Psychol"},{"doi-asserted-by":"publisher","unstructured":"Egly R, Driver J, Rafal R (1994) Shifting visual attention between objects and locations: evidence from normal and parietal lesion subjects. J Exp Psychol Gen 123(2):161\u2013177. https:\/\/doi.org\/10.1037\/\/0096-3445.123.2.161","key":"20736_CR23","DOI":"10.1037\/\/0096-3445.123.2.161"},{"issue":"1","key":"20736_CR24","doi-asserted-by":"publisher","first-page":"52","DOI":"10.3758\/BF03194557","volume":"64","author":"D Lamy","year":"2002","unstructured":"Lamy D, Egeth H (2002) Object-based selection: The role of attentional shifts. Percept Psychophys 64(1):52\u201366. https:\/\/doi.org\/10.3758\/BF03194557","journal-title":"Percept Psychophys"},{"key":"20736_CR25","doi-asserted-by":"publisher","first-page":"3","DOI":"10.1080\/00335558008248231","volume":"32","author":"M Posner","year":"1980","unstructured":"Posner M (1980) Orienting of attention. Q J Exp Psychol 32:3\u201325. https:\/\/doi.org\/10.1080\/00335558008248231","journal-title":"Q J Exp Psychol"},{"issue":"1","key":"20736_CR26","doi-asserted-by":"publisher","first-page":"88","DOI":"10.1186\/s40537-024-00944-3","volume":"11","author":"G Mostafa","year":"2024","unstructured":"Mostafa G, Mahmoud H, Abd El-Hafeez T, ElAraby ME (2024) Feature reduction for hepatocellular carcinoma prediction using machine learning algorithms. J Big Data 11(1):88","journal-title":"J Big Data"},{"issue":"3","key":"20736_CR27","first-page":"697","volume":"1","author":"MR Girgis","year":"2007","unstructured":"Girgis MR, Mahmoud TM, Abd-El-Hafeez T (2007) An approach to image extraction and accurate skin detection from web pages. Int J Comput Inf Eng 1(3):697\u2013705","journal-title":"Int J Comput Inf Eng"},{"issue":"2","key":"20736_CR28","first-page":"2838","volume":"9","author":"T Abd El-Hafeez","year":"2010","unstructured":"Abd El-Hafeez T (2010) A new system for extracting and detecting skin color regions from pdf documents. Int J Comput Sci Eng (IJCSE) 9(2):2838\u20132846","journal-title":"Int J Comput Sci Eng (IJCSE)"},{"unstructured":"El-Sayed MA, Hafeez TA-E (2012) New edge detection technique based on the shannon entropy in gray level images. arXiv preprint arXiv:1211.2502","key":"20736_CR29"},{"doi-asserted-by":"publisher","unstructured":"Naoum A, Nothman J, Curran J (2019) Article segmentation in digitised newspapers with a 2D Markov model. In: Proceedings of the 2019 international conference on document analysis and recognition, pp 1007\u20131014. https:\/\/doi.org\/10.1109\/ICDAR.2019.00165","key":"20736_CR30","DOI":"10.1109\/ICDAR.2019.00165"},{"doi-asserted-by":"publisher","unstructured":"Gu Z, Meng C, Wang K, Lan, J, Wang W, Gu M, Zhang L (2022) XYLayoutLM: Towards layout-aware multimodal networks for visually-rich document understanding. https:\/\/doi.org\/10.48550\/arxiv.2203.06947","key":"20736_CR31","DOI":"10.48550\/arxiv.2203.06947"},{"key":"20736_CR32","first-page":"12077","volume":"34","author":"E Xie","year":"2021","unstructured":"Xie E, Wang W, Yu Z, Anandkumar A, Alvarez JM, Luo P (2021) Segformer: Simple and efficient design for semantic segmentation with transformers. Adv Neural Inf Process Syst 34:12077\u201312090","journal-title":"Adv Neural Inf Process Syst"},{"doi-asserted-by":"publisher","unstructured":"Prasad A, D\u00e9jean H, Meunier J-L (2019) Versatile layout understanding via conjugate graph. In: Proceedings of the 2019 international conference on document analysis and recognition, pp 287\u2013294. https:\/\/doi.org\/10.1109\/ICDAR.2019.00054","key":"20736_CR33","DOI":"10.1109\/ICDAR.2019.00054"},{"doi-asserted-by":"publisher","unstructured":"Xu Y, Li M, Cui L, Huang S, Wei F, Zhou M (2020) LayoutLM: Pre-training of text and layout for document image understanding. In: Proceedings of the 26th ACM SIGKDD international conference on knowledge discovery and data mining, pp 1192\u20131200. https:\/\/doi.org\/10.1145\/3394486.3403172","key":"20736_CR34","DOI":"10.1145\/3394486.3403172"},{"key":"20736_CR35","doi-asserted-by":"publisher","first-page":"359","DOI":"10.1007\/s10791-005-6991-7","volume":"8","author":"A Trotman","year":"2005","unstructured":"Trotman A (2005) Learning to rank. Inf Retr 8:359\u2013381. https:\/\/doi.org\/10.1007\/s10791-005-6991-7","journal-title":"Learning to rank. Inf Retr"},{"key":"20736_CR36","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-642-14267-3","author":"T-Y Liu","year":"2011","unstructured":"Liu T-Y (2011) Learning to Rank for Information Retrieval. Springer, Berlin, Heidelberg. https:\/\/doi.org\/10.1007\/978-3-642-14267-3","journal-title":"Springer, Berlin, Heidelberg."},{"doi-asserted-by":"publisher","unstructured":"Howard A, Sandler M, Chen B, Wang W, Chen L, Tan M, Chu G, Vasudevan V, Zhu Y, Pang R, Adam H, Le Q (2019) Searching for MobileNetV3. In: Proceedings of the 2019 IEEE\/CVF international conference on computer vision, pp 1314\u20131324. https:\/\/doi.org\/10.1109\/ICCV.2019.00140","key":"20736_CR37","DOI":"10.1109\/ICCV.2019.00140"},{"doi-asserted-by":"publisher","unstructured":"Kingma DP, Ba J (2014) Adam: A method for stochastic optimization. https:\/\/doi.org\/10.48550\/arxiv.1412.6980","key":"20736_CR38","DOI":"10.48550\/arxiv.1412.6980"},{"doi-asserted-by":"publisher","unstructured":"Mukherjee K, Khare A, Verma A (2019) A simple dynamic learning rate tuning algorithm for automated training of DNNs. https:\/\/doi.org\/10.48550\/arxiv.1910.11605","key":"20736_CR39","DOI":"10.48550\/arxiv.1910.11605"},{"doi-asserted-by":"publisher","unstructured":"Liao M, Cloud H, Shenzhen, (2023) Real- time scene text detection with differentiable binarization and adaptive scale fusion. IEEE Trans Pattern Anal Mach Intell 919\u2013931. https:\/\/doi.org\/10.1109\/TPAMI.2022.3155612","key":"20736_CR40","DOI":"10.1109\/TPAMI.2022.3155612"},{"doi-asserted-by":"publisher","unstructured":"Du Y, Chen Z, Jia C, Yin X, Zheng T, Li C, Du Y, Jiang Y-G (2022) Svtr: Scene text recognition with a single visual model. In: Proceedings of the 31st international joint conference on artificial intelligence, pp 884\u2013890. https:\/\/doi.org\/10.24963\/ijcai.2022\/124","key":"20736_CR41","DOI":"10.24963\/ijcai.2022\/124"},{"doi-asserted-by":"publisher","unstructured":"Kumar R, Vassilvitskii S (2010) Generalized distances between rankings. In: Proceedings of the 19th international conference on world wide web, pp 571\u2013580. https:\/\/doi.org\/10.1145\/1772690.1772749","key":"20736_CR42","DOI":"10.1145\/1772690.1772749"},{"issue":"1\u20132","key":"20736_CR43","doi-asserted-by":"publisher","first-page":"81","DOI":"10.1093\/biomet\/30.1-2.81","volume":"30","author":"MG Kendall","year":"1938","unstructured":"Kendall MG (1938) A new measure of rank correlation. Biometrika 30(1\u20132):81\u201393. https:\/\/doi.org\/10.1093\/biomet\/30.1-2.81","journal-title":"Biometrika"},{"key":"20736_CR44","doi-asserted-by":"publisher","first-page":"30174","DOI":"10.1109\/ACCESS.2018.2840218","volume":"6","author":"H Yang","year":"2018","unstructured":"Yang H, Jin L, Huang W, Yang Z, Lai S, Sun J (2018) Dense and tight detection of Chinese characters in historical documents: Datasets and a recognition guided detector. IEEE Access 6:30174\u201330183. https:\/\/doi.org\/10.1109\/ACCESS.2018.2840218","journal-title":"IEEE Access"},{"doi-asserted-by":"publisher","unstructured":"Ma W, Zhang H, Jin L, Wu S, Wang J, Wang Y (2020) Joint layout analysis, character detection and recognition for historical document digitization. In: Proceedings of the 17th international conference on frontiers in handwriting recognition, pp 31\u201336. https:\/\/doi.org\/10.48550\/arxiv.2007.06890","key":"20736_CR45","DOI":"10.48550\/arxiv.2007.06890"},{"doi-asserted-by":"publisher","unstructured":"Pedregosa F, Varoquaux G, Gramfort A, Michel V, Thirion B, Grisel O, Blondel M, Prettenhofer P, Weiss R, Dubourg V, Vanderplas J, Passos A, Cournapeau D, Brucher M, Perrot M, Duchesnay E (2011) Scikit-learn: Machine learning in Python. J Mach Learn Res 12:2825\u20132830. https:\/\/doi.org\/10.48550\/arXiv.1201.0490","key":"20736_CR46","DOI":"10.48550\/arXiv.1201.0490"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-025-20736-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-025-20736-y\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-025-20736-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,26]],"date-time":"2025-09-26T15:12:37Z","timestamp":1758899557000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-025-20736-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,3,15]]},"references-count":46,"journal-issue":{"issue":"32","published-online":{"date-parts":[[2025,9]]}},"alternative-id":["20736"],"URL":"https:\/\/doi.org\/10.1007\/s11042-025-20736-y","relation":{},"ISSN":["1573-7721"],"issn-type":[{"type":"electronic","value":"1573-7721"}],"subject":[],"published":{"date-parts":[[2025,3,15]]},"assertion":[{"value":"16 April 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 December 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 March 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"15 March 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest\/Competing interests"}}]}}