{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,25]],"date-time":"2026-03-25T06:04:50Z","timestamp":1774418690741,"version":"3.50.1"},"reference-count":94,"publisher":"Springer Science and Business Media LLC","issue":"18","license":[{"start":{"date-parts":[[2023,11,30]],"date-time":"2023-11-30T00:00:00Z","timestamp":1701302400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2023,11,30]],"date-time":"2023-11-30T00:00:00Z","timestamp":1701302400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-023-17671-1","type":"journal-article","created":{"date-parts":[[2023,11,30]],"date-time":"2023-11-30T08:02:27Z","timestamp":1701331347000},"page":"55773-55810","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":7,"title":["Natural scene text localization and detection using MSER and its variants: a comprehensive survey"],"prefix":"10.1007","volume":"83","author":[{"given":"Kalpita","family":"Dutta","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ritesh","family":"Sarkhel","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mahantapas","family":"Kundu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mita","family":"Nasipuri","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-2426-9915","authenticated-orcid":false,"given":"Nibaran","family":"Das","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2023,11,30]]},"reference":[{"key":"17671_CR1","doi-asserted-by":"publisher","unstructured":"Ajay B, Naveena C (2019) A mechanism for detection of text in images using dwt and mser. Integr Intell Comput Commun Secur 669\u20136760. https:\/\/doi.org\/10.1007\/978-981-10-8797-4_68","DOI":"10.1007\/978-981-10-8797-4_68"},{"issue":"23","key":"17671_CR2","doi-asserted-by":"publisher","first-page":"34047","DOI":"10.1007\/s11042-022-13179-2","volume":"81","author":"A Akoushideh","year":"2022","unstructured":"Akoushideh A, Rasoulnejad SMF, Shahbahrami A (2022) Text localization in digital images using a hybrid method. Multimed Tools Appl 81(23):34047\u201334066. https:\/\/doi.org\/10.1007\/s11042-022-13179-2","journal-title":"Multimed Tools Appl"},{"key":"17671_CR3","unstructured":"Ali H (2022) Leveraging machine learning for less developed languages: progress on urdu text detection. arXiv:2209.14022"},{"issue":"1","key":"17671_CR4","first-page":"71","volume":"39","author":"A Awoke","year":"2021","unstructured":"Awoke A, Tekeba M (2021) Ethiopic and latin multilingual text detection from images using hybrid techniques. Zede J 39(1):71\u201380","journal-title":"Zede J"},{"key":"17671_CR5","doi-asserted-by":"publisher","unstructured":"Baek Y, Lee B, Han D et\u00a0al (2019) Character region awareness for text detection. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 9365\u20139374. https:\/\/doi.org\/10.1109\/CVPR.2019.00959","DOI":"10.1109\/CVPR.2019.00959"},{"key":"17671_CR6","doi-asserted-by":"publisher","unstructured":"Bartz C, Yang H, Meinel C (2018) See: towards semi-supervised end-to-end scene text recognition. In: Proceedings of the AAAI conference on artificial intelligence. https:\/\/doi.org\/10.1609\/aaai.v32i1.12242","DOI":"10.1609\/aaai.v32i1.12242"},{"key":"17671_CR7","doi-asserted-by":"publisher","unstructured":"Busta M, Neumann L, Matas J (2017) Deep textspotter: an end-to-end trainable scene text localization and recognition framework. In: Proceedings of the IEEE international conference on computer vision, pp 2204\u20132212. https:\/\/doi.org\/10.1109\/ICCV.2017.242","DOI":"10.1109\/ICCV.2017.242"},{"key":"17671_CR8","doi-asserted-by":"publisher","unstructured":"Chaitra Y, Dinesh R (2022) An impact of radon transforms and filtering techniques for text localization in natural scene text images. In: ICT with intelligent applications: proceedings of ICTIS 2021. Springer, vol 1, pp 563\u2013573. https:\/\/doi.org\/10.1007\/978-981-16-4177-0_55","DOI":"10.1007\/978-981-16-4177-0_55"},{"key":"17671_CR9","doi-asserted-by":"publisher","unstructured":"Chen H, Tsai SS, Schroth G et\u00a0al (2011) Robust text detection in natural images with edge-enhanced maximally stable extremal regions. In: 2011 18th IEEE international conference on image processing. IEEE, pp 2609\u20132612. https:\/\/doi.org\/10.1109\/ICIP.2011.6116200","DOI":"10.1109\/ICIP.2011.6116200"},{"key":"17671_CR10","doi-asserted-by":"publisher","unstructured":"Chen X, Jin L, Zhu Y et\u00a0al (2021) Text recognition in the wild: a survey. ACM, New York, vol\u00a054, pp 1\u201335. https:\/\/doi.org\/10.1145\/3440756","DOI":"10.1145\/3440756"},{"key":"17671_CR11","doi-asserted-by":"crossref","unstructured":"Cho H, Sung M, Jun B (2016) Canny text detector: fast and robust scene text localization algorithm. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3566\u20133573","DOI":"10.1109\/CVPR.2016.388"},{"key":"17671_CR12","doi-asserted-by":"publisher","unstructured":"Choudhary S, Singh NK, Chichadwani S (2018) Text detection and recognition from scene images using mser and cnn. In: 2018 second international conference on advances in electronics, computers and communications (ICAECC). IEEE, pp 1\u20134.https:\/\/doi.org\/10.1109\/ICAECC.2018.8479419","DOI":"10.1109\/ICAECC.2018.8479419"},{"key":"17671_CR13","doi-asserted-by":"publisher","DOI":"10.2307\/2583667","author":"TH Cormen","year":"2009","unstructured":"Cormen TH, Leiserson CE, Rivest RL et al (2009) Introduction to algorithms. MIT Press. https:\/\/doi.org\/10.2307\/2583667","journal-title":"MIT Press"},{"key":"17671_CR14","doi-asserted-by":"publisher","unstructured":"Das S, Chattopadhyay S, Prasad R et\u00a0al (2022) Text region identification from natural scene images using semi-supervised mser method. In: Proceedings of 2nd international conference on mathematical modeling and computational science: ICMMCS 2021. Springer, pp 401\u2013408. https:\/\/doi.org\/10.1007\/978-981-19-0182-9_40","DOI":"10.1007\/978-981-19-0182-9_40"},{"key":"17671_CR15","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1155\/2020\/7067251","volume":"2020","author":"J Diaz-Escobar","year":"2020","unstructured":"Diaz-Escobar J, Kober V (2020) Natural scene text detection and segmentation using phase-based regions and character retrieval. Math Probl Eng 2020:1\u201317. https:\/\/doi.org\/10.1155\/2020\/7067251","journal-title":"Math Probl Eng"},{"key":"17671_CR16","doi-asserted-by":"publisher","unstructured":"El\u00a0Abbadi NK et\u00a0al (2023) Scene text detection and recognition by using multi-level features extractions based on you only once version five (yolov5) and maximally stable extremal regions (msers) with optical character recognition (ocr). Al-Salam J Eng Technol 2(1):13\u201327. https:\/\/doi.org\/10.55145\/ajest.2023.01.01.002","DOI":"10.55145\/ajest.2023.01.01.002"},{"key":"17671_CR17","doi-asserted-by":"publisher","unstructured":"Ghosh J, Talukdar AK, Sarma KK (2023) A light-weight natural scene text detection and recognition system. Multimed Tools Appl 1\u201333. https:\/\/doi.org\/10.1007\/s11042-023-15696-0","DOI":"10.1007\/s11042-023-15696-0"},{"key":"17671_CR18","doi-asserted-by":"publisher","unstructured":"Goud DS, Vigneshwari M, Aparna P et\u00a0al (2022) Text localization and recognition from natural scene images using ai. In: 2022 International conference on automation, computing and renewable systems (ICACRS). IEEE, pp 1153\u20131158. https:\/\/doi.org\/10.1109\/ICACRS55517.2022.10029220","DOI":"10.1109\/ICACRS55517.2022.10029220"},{"key":"17671_CR19","doi-asserted-by":"publisher","first-page":"10821","DOI":"10.1007\/s11042-018-6613-1","volume":"8","author":"N Gupta","year":"2019","unstructured":"Gupta N, Jalal AS (2019) A robust model for salient text detection in natural scene images using mser feature detector and grabcut. Multimed Tools Appl 8:10821\u201310835. https:\/\/doi.org\/10.1007\/s11042-018-6613-1","journal-title":"Multimed Tools Appl"},{"issue":"6","key":"17671_CR20","doi-asserted-by":"publisher","first-page":"1397","DOI":"10.1109\/TPAMI.2012.213","volume":"35","author":"K He","year":"2012","unstructured":"He K, Sun J, Tang X (2012) Guided image filtering. IEEE Trans Pattern Anal Mach Intell 35(6):1397\u20131409. https:\/\/doi.org\/10.1109\/TPAMI.2012.213","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"17671_CR21","doi-asserted-by":"publisher","unstructured":"He M, Liao M, Yang Z et\u00a0al (2021) Most: a multi-oriented scene text detector with localization refinement. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 8813\u20138822. https:\/\/doi.org\/10.1109\/CVPR46437.2021.00870","DOI":"10.1109\/CVPR46437.2021.00870"},{"issue":"6","key":"17671_CR22","doi-asserted-by":"publisher","first-page":"2529","DOI":"10.1109\/TIP.2016.2547588","volume":"25","author":"T He","year":"2016","unstructured":"He T, Huang W, Qiao Y et al (2016) Text-attentional convolutional neural network for scene text detection. IEEE Trans Image Process 25(6):2529\u20132541. https:\/\/doi.org\/10.1109\/TIP.2016.2547588","journal-title":"IEEE Trans Image Process"},{"key":"17671_CR23","doi-asserted-by":"publisher","unstructured":"He W, Zhang XY, Yin F et\u00a0al (2017) Deep direct regression for multi-oriented scene text detection. In: Proceedings of the IEEE international conference on computer vision, pp 745\u2013753. https:\/\/doi.org\/10.1109\/ICCV.2017.87","DOI":"10.1109\/ICCV.2017.87"},{"key":"17671_CR24","doi-asserted-by":"publisher","unstructured":"Huang W, Qiao Y, Tang X (2014) Robust scene text detection with convolution neural network induced mser trees. In: European conference on computer vision. Springer, pp 497\u2013511. https:\/\/doi.org\/10.1007\/978-3-319-10593-2_33","DOI":"10.1007\/978-3-319-10593-2_33"},{"key":"17671_CR25","doi-asserted-by":"publisher","unstructured":"Islam MR, Mondal C, Azam MK et\u00a0al (2016) Text detection and recognition using enhanced mser detection and a novel ocr technique. In: 2016 5th international conference on informatics, electronics and vision (ICIEV). IEEE, pp 15\u201320. https:\/\/doi.org\/10.1109\/ICIEV.2016.7760054","DOI":"10.1109\/ICIEV.2016.7760054"},{"key":"17671_CR26","doi-asserted-by":"publisher","unstructured":"Islam R, Islam R, Talukder KH (2020) An enhanced mser pruning algorithm for detection and localization of bangla texts from scene images. Int Arab J Inf Technol 17(3):375\u2013385. https:\/\/doi.org\/10.34028\/iajit\/17\/3\/11","DOI":"10.34028\/iajit\/17\/3\/11"},{"key":"17671_CR27","doi-asserted-by":"crossref","unstructured":"Jiang Y, Zhu X, Wang X et\u00a0al (2017) R2cnn: rotational region cnn for orientation robust scene text detection. arXiv:1706.09579","DOI":"10.1109\/ICPR.2018.8545598"},{"issue":"1","key":"17671_CR28","doi-asserted-by":"publisher","first-page":"78","DOI":"10.4218\/etrij.11.1510.0029","volume":"33","author":"J Jung","year":"2011","unstructured":"Jung J, Lee S, Cho MS et al (2011) Touch tt: scene text extractor using touchscreen interface. ETRI J 33(1):78\u201388. https:\/\/doi.org\/10.4218\/etrij.11.1510.0029","journal-title":"ETRI J"},{"issue":"5","key":"17671_CR29","doi-asserted-by":"publisher","first-page":"977","DOI":"10.1016\/j.patcog.2003.10.012","volume":"37","author":"K Jung","year":"2004","unstructured":"Jung K, Kim KI, Jain AK (2004) Text information extraction in images and video: a survey. Pattern Recognit 37(5):977\u2013997. https:\/\/doi.org\/10.1016\/j.patcog.2003.10.012","journal-title":"Pattern Recognit"},{"key":"17671_CR30","doi-asserted-by":"publisher","unstructured":"Karaoglu S, Fernando B, Tr\u00e9meau A (2010) A novel algorithm for text detection and localization in natural scene images. In: 2010 international conference on digital image computing: techniques and applications. IEEE, pp 635\u2013642. https:\/\/doi.org\/10.1109\/DICTA.2010.115","DOI":"10.1109\/DICTA.2010.115"},{"key":"17671_CR31","doi-asserted-by":"publisher","unstructured":"Karatzas D, Shafait F, Uchida S et\u00a0al (2013) Icdar 2013 robust reading competition. In: 2013 12th international conference on document analysis and recognition. IEEE, pp 1484\u20131493. https:\/\/doi.org\/10.1109\/ICDAR.2013.221","DOI":"10.1109\/ICDAR.2013.221"},{"key":"17671_CR32","doi-asserted-by":"publisher","unstructured":"Karatzas D, Gomez-Bigorda L, Nicolaou A et\u00a0al (2015) Icdar 2015 competition on robust reading. In: 2015 13th international conference on document analysis and recognition (ICDAR). IEEE, pp 1156\u20131160. https:\/\/doi.org\/10.1109\/ICDAR.2015.7333942","DOI":"10.1109\/ICDAR.2015.7333942"},{"key":"17671_CR33","doi-asserted-by":"publisher","first-page":"128","DOI":"10.1016\/j.patcog.2016.01.008","volume":"54","author":"V Khare","year":"2016","unstructured":"Khare V, Shivakumara P, Raveendran P et al (2016) A blind deconvolution model for scene text detection and recognition in video. Pattern Recognit 54:128\u2013148. https:\/\/doi.org\/10.1016\/j.patcog.2016.01.008","journal-title":"Pattern Recognit"},{"key":"17671_CR34","unstructured":"L1kw1d (2013) Borndigitaltext. Accessed 18 June 2023"},{"issue":"7","key":"17671_CR35","doi-asserted-by":"publisher","first-page":"10595","DOI":"10.1007\/s11042-022-13690-6","volume":"82","author":"G Larbi","year":"2023","unstructured":"Larbi G (2023) Two-step text detection framework in natural scenes based on pseudo-zernike moments and cnn. Multimed Tools Appl 82(7):10595\u201310616. https:\/\/doi.org\/10.1007\/s11042-022-13690-6","journal-title":"Multimed Tools Appl"},{"key":"17671_CR36","doi-asserted-by":"publisher","unstructured":"Lee CY, Baek Y, Lee H (2019) Tedeval: a fair evaluation metric for scene text detectors. In: 2019 international conference on document analysis and recognition workshops (ICDARW). IEEE, pp 14\u201317. https:\/\/doi.org\/10.1109\/icdarw.2019.60125","DOI":"10.1109\/icdarw.2019.60125"},{"issue":"1","key":"17671_CR37","doi-asserted-by":"publisher","first-page":"1","DOI":"10.4018\/jdm.322086","volume":"34","author":"R Li","year":"2023","unstructured":"Li R, Chen S, Zhao F et al (2023) Text detection model for historical documents using cnn and mser. J Database Manag (JDM) 34(1):1\u201323. https:\/\/doi.org\/10.4018\/jdm.322086","journal-title":"J Database Manag (JDM)"},{"key":"17671_CR38","unstructured":"Li Y, Lu H (2012) Scene text detection via stroke width. In: Proceedings of the 21st international conference on pattern recognition (ICPR2012). IEEE, pp 681\u2013684"},{"key":"17671_CR39","doi-asserted-by":"publisher","unstructured":"Li Y, Shen C, Jia W et\u00a0al (2013) Leveraging surrounding context for scene text detection. In: 2013 IEEE international conference on image processing. IEEE, pp 2264\u20132268. https:\/\/doi.org\/10.1109\/ICIP.2013.6738467","DOI":"10.1109\/ICIP.2013.6738467"},{"issue":"4","key":"17671_CR40","doi-asserted-by":"publisher","first-page":"1666","DOI":"10.1109\/TIP.2014.2302896","volume":"23","author":"Y Li","year":"2014","unstructured":"Li Y, Jia W, Shen C et al (2014) Characterness: an indicator of text in the wild. IEEE Trans Image Process 23(4):1666\u20131677. https:\/\/doi.org\/10.1109\/TIP.2014.2302896","journal-title":"IEEE Trans Image Process"},{"issue":"2","key":"17671_CR41","doi-asserted-by":"publisher","first-page":"84","DOI":"10.1007\/s10032-004-0138-z","volume":"7","author":"J Liang","year":"2005","unstructured":"Liang J, Doermann D, Li H (2005) Camera-based analysis of text and documents: a survey. Int J Doc Anal Recognit (IJDAR) 7(2):84\u2013104. https:\/\/doi.org\/10.1007\/s10032-004-0138-z","journal-title":"Int J Doc Anal Recognit (IJDAR)"},{"key":"17671_CR42","doi-asserted-by":"publisher","unstructured":"Liao M, Shi B, Bai X et\u00a0al (2017) Textboxes: a fast text detector with a single deep neural network. In: Proceedings of the AAAI conference on artificial intelligence. https:\/\/doi.org\/10.1609\/aaai.v31i1.11196","DOI":"10.1609\/aaai.v31i1.11196"},{"issue":"8","key":"17671_CR43","doi-asserted-by":"publisher","first-page":"3676","DOI":"10.1109\/TIP.2018.2825107","volume":"27","author":"M Liao","year":"2018","unstructured":"Liao M, Shi B, Bai X (2018) Textboxes++: a single-shot oriented scene text detector. IEEE Transac Image Process 27(8):3676\u20133690. https:\/\/doi.org\/10.1109\/TIP.2018.2825107","journal-title":"IEEE Transac Image Process"},{"key":"17671_CR44","doi-asserted-by":"publisher","unstructured":"Liu Y, Jin L (2017) Deep matching prior network: toward tighter multi-oriented text detection. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 1962\u20131969. https:\/\/doi.org\/10.1109\/CVPR.2017.368","DOI":"10.1109\/CVPR.2017.368"},{"key":"17671_CR45","doi-asserted-by":"publisher","unstructured":"Liu Y, Jin L, Xie Z et\u00a0al (2019) Tightness-aware evaluation protocol for scene text detection. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 9612\u20139620. https:\/\/doi.org\/10.1109\/CVPR.2019.00984","DOI":"10.1109\/CVPR.2019.00984"},{"issue":"1","key":"17671_CR46","doi-asserted-by":"publisher","first-page":"161","DOI":"10.1007\/s11263-020-01369-0","volume":"129","author":"S Long","year":"2021","unstructured":"Long S, He X, Yao C (2021) Scene text detection and recognition: the deep learning era. Int J Comput Vis 129(1):161\u2013184. https:\/\/doi.org\/10.1007\/s11263-020-01369-0","journal-title":"Int J Comput Vis"},{"issue":"2\u20133","key":"17671_CR47","doi-asserted-by":"publisher","first-page":"105","DOI":"10.1007\/s10032-004-0134-3","volume":"7","author":"SM Lucas","year":"2005","unstructured":"Lucas SM, Panaretos A, Sosa L et al (2005) Icdar 2003 robust reading competitions: entries, results, and future directions. Int J Doc Anal Recognit (IJDAR) 7(2\u20133):105\u2013122. https:\/\/doi.org\/10.1007\/s10032-004-0134-3","journal-title":"Int J Doc Anal Recognit (IJDAR)"},{"key":"17671_CR48","doi-asserted-by":"publisher","unstructured":"Lundgren A, Castro D, Lima E et\u00a0al (2019) Octshufflemlt: a compact octave based neural network for end-to-end multilingual text detection and recognition. In: 2019 international conference on document analysis and recognition workshops (ICDARW). IEEE, pp 37\u201342. https:\/\/doi.org\/10.1109\/ICDARW.2019.30062","DOI":"10.1109\/ICDARW.2019.30062"},{"key":"17671_CR49","doi-asserted-by":"publisher","unstructured":"Lyu P, Liao M, Yao C et\u00a0al (2018a) Mask textspotter: an end-to-end trainable neural network for spotting text with arbitrary shapes. In: Proceedings of the European conference on computer vision (ECCV), pp 67\u201383. https:\/\/doi.org\/10.1109\/TPAMI.2019.2937086","DOI":"10.1109\/TPAMI.2019.2937086"},{"key":"17671_CR50","doi-asserted-by":"publisher","unstructured":"Lyu P, Yao C, Wu W et\u00a0al (2018b) Multi-oriented scene text detection via corner localization and region segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 7553\u20137563. https:\/\/doi.org\/10.1109\/CVPR.2018.00788","DOI":"10.1109\/CVPR.2018.00788"},{"issue":"11","key":"17671_CR51","doi-asserted-by":"publisher","first-page":"3111","DOI":"10.1109\/TMM.2018.2818020","volume":"20","author":"J Ma","year":"2018","unstructured":"Ma J, Shao W, Ye H et al (2018) Arbitrary-oriented scene text detection via rotation proposals. IEEE Trans Multimed 20(11):3111\u20133122. https:\/\/doi.org\/10.1109\/TMM.2018.2818020","journal-title":"IEEE Trans Multimed"},{"key":"17671_CR52","doi-asserted-by":"publisher","unstructured":"Mansouri S, Zrigui S, Zrigui M et\u00a0al (2021) Text detection in Arabic news video based on mser and retinanet. In: 2021 IEEE\/ACS 18th international conference on computer systems and applications (AICCSA). IEEE, pp 1\u20137. https:\/\/doi.org\/10.1109\/AICCSA53542.2021.9686930","DOI":"10.1109\/AICCSA53542.2021.9686930"},{"issue":"10","key":"17671_CR53","doi-asserted-by":"publisher","first-page":"761","DOI":"10.1016\/j.imavis.2004.02.006","volume":"22","author":"J Matas","year":"2004","unstructured":"Matas J, Chum O, Urban M et al (2004) Robust wide-baseline stereo from maximally stable extremal regions. Image Vis Comput 22(10):761\u2013767. https:\/\/doi.org\/10.1016\/j.imavis.2004.02.006","journal-title":"Image Vis Comput"},{"key":"17671_CR54","doi-asserted-by":"crossref","unstructured":"Meetei LS, Singh TD, Bandyopadhyay S (2019) Extraction and identification of manipuri and mizo texts from scene and document images. In: International conference on pattern recognition and machine intelligence. Springer, pp 405\u2013414","DOI":"10.1007\/978-3-030-34869-4_44"},{"key":"17671_CR55","doi-asserted-by":"publisher","first-page":"27137","DOI":"10.1007\/s11042-020-09318-2","volume":"79","author":"F Naiemi","year":"2020","unstructured":"Naiemi F, Ghods V, Khalesi H (2020) Scene text detection using enhanced extremal region and convolutional neural network. Multimed Tools Appl 79:27137\u201327159. https:\/\/doi.org\/10.1007\/s11042-020-09318-2","journal-title":"Multimed Tools Appl"},{"key":"17671_CR56","doi-asserted-by":"publisher","DOI":"10.1016\/j.eswa.2020.114549","volume":"170","author":"F Naiemi","year":"2021","unstructured":"Naiemi F, Ghods V, Khalesi H (2021) A novel pipeline framework for multi oriented scene text image detection and recognition. Expert Syst Appl 170:114549. https:\/\/doi.org\/10.1016\/j.eswa.2020.114549","journal-title":"Expert Syst Appl"},{"key":"17671_CR57","doi-asserted-by":"publisher","unstructured":"Nayef N, Yin F, Bizid I et\u00a0al (2017) Icdar2017 robust reading challenge on multi-lingual scene text detection and script identification-rrc-mlt. In: 2017 14th IAPR international conference on document analysis and recognition (ICDAR). IEEE, pp 1454\u20131459. https:\/\/doi.org\/10.1109\/ICDAR.2017.237","DOI":"10.1109\/ICDAR.2017.237"},{"key":"17671_CR58","doi-asserted-by":"publisher","unstructured":"Nayef N, Patel Y, Busta M et\u00a0al (2019) Icdar2019 robust reading challenge on multi-lingual scene text detection and recognition-rrc-mlt-2019. In: 2019 International conference on document analysis and recognition (ICDAR). IEEE, pp 1582\u20131587. https:\/\/doi.org\/10.1109\/ICDAR.2019.00254","DOI":"10.1109\/ICDAR.2019.00254"},{"key":"17671_CR59","doi-asserted-by":"publisher","unstructured":"Panda S, Ash S, Chakraborty N et\u00a0al (2020) Parameter tuning in mser for text localization in multi-lingual camera-captured scene text images. In: Computational intelligence in pattern recognition: proceedings of CIPR 2019. Springer, pp 999\u20131009. https:\/\/doi.org\/10.1007\/978-981-13-9042-5_86","DOI":"10.1007\/978-981-13-9042-5_86"},{"key":"17671_CR60","doi-asserted-by":"publisher","unstructured":"Qin L, Shivakumara P, Lu T et\u00a0al (2016) Video scene text frames categorization for text detection and recognition. In: 2016 23rd international conference on pattern recognition (ICPR). IEEE, pp 3886\u20133891. https:\/\/doi.org\/10.1109\/ICPR.2016.7900241","DOI":"10.1109\/ICPR.2016.7900241"},{"key":"17671_CR61","doi-asserted-by":"crossref","unstructured":"Shahab A, Shafait F, Dengel A (2011) Icdar 2011 robust reading competition challenge 2: Reading text in scene images. In: 2011 international conference on document analysis and recognition. IEEE, pp 1491\u20131496","DOI":"10.1109\/ICDAR.2011.296"},{"issue":"9","key":"17671_CR62","doi-asserted-by":"publisher","first-page":"2853","DOI":"10.1016\/j.patcog.2014.03.023","volume":"47","author":"C Shi","year":"2014","unstructured":"Shi C, Wang C, Xiao B et al (2014) End-to-end scene text recognition using tree-structured models. Pattern Recognit 47(9):2853\u20132866. https:\/\/doi.org\/10.1016\/j.patcog.2014.03.023","journal-title":"Pattern Recognit"},{"issue":"2","key":"17671_CR63","doi-asserted-by":"publisher","first-page":"412","DOI":"10.1109\/TPAMI.2010.166","volume":"33","author":"P Shivakumara","year":"2010","unstructured":"Shivakumara P, Phan TQ, Tan CL (2010) A laplacian approach to multi-oriented text detection in video. IEEE Trans Pattern Anal Mach Intell 33(2):412\u2013419. https:\/\/doi.org\/10.1109\/TPAMI.2010.166","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"issue":"8","key":"17671_CR64","doi-asserted-by":"publisher","first-page":"1227","DOI":"10.1109\/TCSVT.2012.2198129","volume":"22","author":"P Shivakumara","year":"2012","unstructured":"Shivakumara P, Sreedhar RP, Phan TQ et al (2012) Multioriented video scene text detection through bayesian classification and boundary growing. IEEE Trans Circ Syst Vid Technol 22(8):1227\u20131235. https:\/\/doi.org\/10.1109\/TCSVT.2012.2198129","journal-title":"IEEE Trans Circ Syst Vid Technol"},{"key":"17671_CR65","doi-asserted-by":"publisher","unstructured":"Soni R, Kumar B, Chand S (2017) Text detection and localization in natural scene images using mser and fast guided filter. In: 2017 fourth international conference on image information processing (ICIIP). IEEE, pp 1\u20136. https:\/\/doi.org\/10.1109\/ICIIP.2017.8313739","DOI":"10.1109\/ICIIP.2017.8313739"},{"key":"17671_CR66","doi-asserted-by":"publisher","unstructured":"Tabassum A, Dhondse SA (2015) Text detection using mser and stroke width transform. In: 2015 fifth international conference on communication systems and network technologies. IEEE, pp 568\u2013571. https:\/\/doi.org\/10.1109\/CSNT.2015.154","DOI":"10.1109\/CSNT.2015.154"},{"issue":"5","key":"17671_CR67","doi-asserted-by":"publisher","first-page":"11681","DOI":"10.1007\/s10586-017-1448-5","volume":"22","author":"A Thilagavathy","year":"2019","unstructured":"Thilagavathy A, Chilambuchelvan A (2019) Fuzzy based edge enhanced text detection algorithm using mser. Clust Comput 22(5):11681\u201311687. https:\/\/doi.org\/10.1007\/s10586-017-1448-5","journal-title":"Clust Comput"},{"key":"17671_CR68","doi-asserted-by":"publisher","unstructured":"Tian S, Lu S, Su B et\u00a0al (2014) Scene text segmentation with multi-level maximally stable extremal regions. In: 2014 22nd international conference on pattern recognition. IEEE, pp 2703\u20132708. https:\/\/doi.org\/10.1109\/ICPR.2014.467","DOI":"10.1109\/ICPR.2014.467"},{"key":"17671_CR69","doi-asserted-by":"publisher","unstructured":"Tian S, Pan Y, Huang C et\u00a0al (2015) Text flow: a unified text detection system in natural scene images. In: Proceedings of the IEEE international conference on computer vision, pp 4651\u20134659. https:\/\/doi.org\/10.1109\/ICCV.2015.528","DOI":"10.1109\/ICCV.2015.528"},{"key":"17671_CR70","doi-asserted-by":"publisher","DOI":"10.1016\/j.knosys.2022.109040","volume":"250","author":"G Tong","year":"2022","unstructured":"Tong G, Dong M, Sun X et al (2022) Natural scene text detection and recognition based on saturation-incorporated multi-channel mser. Knowl-Based Syst 250:109040. https:\/\/doi.org\/10.1016\/j.knosys.2022.109040","journal-title":"Knowl-Based Syst"},{"key":"17671_CR71","doi-asserted-by":"publisher","unstructured":"Turki H, Halima MB, Alimi AM (2017a) A hybrid method of natural scene text detection using msers masks in hsv space color. In: Ninth international conference on machine vision (ICMV 2016). International Society for Optics and Photonics, p 1034111. https:\/\/doi.org\/10.1117\/12.2268993","DOI":"10.1117\/12.2268993"},{"key":"17671_CR72","doi-asserted-by":"publisher","unstructured":"Turki H, Halima MB, Alimi AM (2017b) Text detection based on mser and cnn features. In: 2017 14th IAPR international conference on document analysis and recognition (ICDAR). IEEE, pp 949\u2013954. https:\/\/doi.org\/10.1109\/ICDAR.2017.159","DOI":"10.1109\/ICDAR.2017.159"},{"key":"17671_CR73","unstructured":"Vishnoitanuj (2020) Handwritten-text. Accessed 18 June 2023"},{"key":"17671_CR74","doi-asserted-by":"publisher","unstructured":"Wan Y, Wang X, Lu D (2019) Research on key technology of Chinese text localization in natural scenes. In: Recent developments in intelligent computing, communication and devices: proceedings of ICCD 2017. Springer, pp 387\u2013397. https:\/\/doi.org\/10.1007\/978-981-10-8944-2_45","DOI":"10.1007\/978-981-10-8944-2_45"},{"key":"17671_CR75","doi-asserted-by":"publisher","unstructured":"Wang K, Babenko B, Belongie S (2011) End-to-end scene text recognition. In: 2011 international conference on computer vision. IEEE, pp 1457\u20131464. https:\/\/doi.org\/10.1109\/ICCV.2011.6126402","DOI":"10.1109\/ICCV.2011.6126402"},{"key":"17671_CR76","unstructured":"Wang T, Wu DJ, Coates A et\u00a0al (2012) End-to-end text recognition with convolutional neural networks. In: Proceedings of the 21st international conference on pattern recognition (ICPR2012). IEEE, pp 3304\u20133308"},{"key":"17671_CR77","doi-asserted-by":"publisher","first-page":"26201","DOI":"10.1007\/s11042-016-4099-2","volume":"76","author":"X Wang","year":"2017","unstructured":"Wang X, Song Y, Zhang Y et al (2017) A hierarchical recursive method for text detection in natural scene images. Multimed Tools Appl 76:26201\u201326223. https:\/\/doi.org\/10.1007\/s11042-016-4099-2","journal-title":"Multimed Tools Appl"},{"key":"17671_CR78","doi-asserted-by":"publisher","first-page":"46","DOI":"10.1016\/j.neucom.2017.12.058","volume":"295","author":"Y Wang","year":"2018","unstructured":"Wang Y, Shi C, Xiao B et al (2018) Crf based text detection for natural scene images using convolutional neural network and context information. Neurocomputing 295:46\u201358. https:\/\/doi.org\/10.1016\/j.neucom.2017.12.058","journal-title":"Neurocomputing"},{"key":"17671_CR79","doi-asserted-by":"publisher","unstructured":"Wang Y, Xie H, Zha ZJ et\u00a0al (2020) Contournet: taking a further step toward accurate arbitrary-shaped scene text detection. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 11753\u201311762. https:\/\/doi.org\/10.1109\/CVPR42600.2020.01177","DOI":"10.1109\/CVPR42600.2020.01177"},{"issue":"4","key":"17671_CR80","doi-asserted-by":"publisher","first-page":"280","DOI":"10.1007\/s10032-006-0014-0","volume":"8","author":"C Wolf","year":"2006","unstructured":"Wolf C, Jolion JM (2006) Object count\/area graphs for the evaluation of object detection and segmentation algorithms. Int J Doc Anal Recognit (IJDAR) 8(4):280\u2013296. https:\/\/doi.org\/10.1007\/s10032-006-0014-0","journal-title":"Int J Doc Anal Recognit (IJDAR)"},{"key":"17671_CR81","doi-asserted-by":"publisher","unstructured":"Xie E, Zang Y, Shao S et\u00a0al (2019) Scene text detection with supervised pyramid context network. In: Proceedings of the AAAI conference on artificial intelligence, pp 9038\u20139045. https:\/\/doi.org\/10.1609\/aaai.v33i01.33019038","DOI":"10.1609\/aaai.v33i01.33019038"},{"key":"17671_CR82","doi-asserted-by":"publisher","unstructured":"Yao C, Bai X, Liu W et\u00a0al (2012) Detecting texts of arbitrary orientations in natural images. In: 2012 IEEE conference on computer vision and pattern recognition. IEEE, pp 1083\u20131090. https:\/\/doi.org\/10.1109\/CVPR.2012.6247787","DOI":"10.1109\/CVPR.2012.6247787"},{"key":"17671_CR83","unstructured":"Yao C, Bai X, Sang N et\u00a0al (2016) Scene text detection via holistic, multi-channel prediction. arXiv:1606.09002"},{"issue":"7","key":"17671_CR84","doi-asserted-by":"publisher","first-page":"1480","DOI":"10.1109\/TPAMI.2014.2366765","volume":"37","author":"Q Ye","year":"2014","unstructured":"Ye Q, Doermann D (2014) Text detection and recognition in imagery: a survey. IEEE Trans Pattern Anal Mach Intell 37(7):1480\u20131500. https:\/\/doi.org\/10.1109\/TPAMI.2014.2366765","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"issue":"6","key":"17671_CR85","doi-asserted-by":"publisher","first-page":"2752","DOI":"10.1109\/TIP.2016.2554321","volume":"25","author":"XC Yin","year":"2016","unstructured":"Yin XC, Zuo ZY, Tian S et al (2016) Text detection, tracking and recognition in video: a comprehensive survey. IEEE Trans Image Process 25(6):2752\u20132773. https:\/\/doi.org\/10.1109\/TIP.2016.2554321","journal-title":"IEEE Trans Image Process"},{"key":"17671_CR86","doi-asserted-by":"publisher","unstructured":"Zhan F, Xue C, Lu S (2019) Ga-dan: geometry-aware domain adaptation network for scene text detection and recognition. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 9105\u20139115. https:\/\/doi.org\/10.1109\/ICCV.2019.00920","DOI":"10.1109\/ICCV.2019.00920"},{"key":"17671_CR87","doi-asserted-by":"publisher","unstructured":"Zhang SX, Zhu X, Hou JB et\u00a0al (2020) Deep relational reasoning graph network for arbitrary shape text detection. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 9699\u20139708. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00972","DOI":"10.1109\/CVPR42600.2020.00972"},{"key":"17671_CR88","doi-asserted-by":"publisher","first-page":"61","DOI":"10.1016\/j.neucom.2018.03.070","volume":"307","author":"X Zhang","year":"2018","unstructured":"Zhang X, Gao X, Tian C (2018) Text detection in natural scene images based on color prior guided mser. Neurocomputing 307:61\u201371. https:\/\/doi.org\/10.1016\/j.neucom.2018.03.070","journal-title":"Neurocomputing"},{"issue":"19","key":"17671_CR89","doi-asserted-by":"publisher","first-page":"29005","DOI":"10.1007\/s11042-021-11101-w","volume":"80","author":"Y Zhang","year":"2021","unstructured":"Zhang Y, Huang Y, Zhao D et al (2021) A scene text detector based on deep feature merging. Multimed Tools Appl 80(19):29005\u201329016. https:\/\/doi.org\/10.1007\/s11042-021-11101-w","journal-title":"Multimed Tools Appl"},{"key":"17671_CR90","doi-asserted-by":"publisher","unstructured":"Zhang Z, Shen W, Yao C et\u00a0al (2015) Symmetry-based text line detection in natural scenes. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2558\u20132567. https:\/\/doi.org\/10.1109\/CVPR.2015.7298871","DOI":"10.1109\/CVPR.2015.7298871"},{"key":"17671_CR91","doi-asserted-by":"publisher","unstructured":"Zhou X, Yao C, Wen H et\u00a0al (2017) East: an efficient and accurate scene text detector. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 5551\u20135560. https:\/\/doi.org\/10.1109\/CVPR.2017.283","DOI":"10.1109\/CVPR.2017.283"},{"key":"17671_CR92","doi-asserted-by":"publisher","unstructured":"Zhu W, Lou J, Chen L et al (2017) Scene text detection via extremal region based double threshold convolutional network classification. PloS One 12(8):e0182227. https:\/\/doi.org\/10.1371\/journal.pone.0182227","DOI":"10.1371\/journal.pone.0182227"},{"issue":"1","key":"17671_CR93","doi-asserted-by":"publisher","first-page":"19","DOI":"10.1007\/s11704-015-4488-0","volume":"10","author":"Y Zhu","year":"2016","unstructured":"Zhu Y, Yao C, Bai X (2016) Scene text detection and recognition: recent advances and future trends. Front Comput Sci 10(1):19\u201336. https:\/\/doi.org\/10.1007\/s11704-015-4488-0","journal-title":"Front Comput Sci"},{"key":"17671_CR94","doi-asserted-by":"publisher","first-page":"62616","DOI":"10.1109\/ACCESS.2019.2916616","volume":"7","author":"LQ Zuo","year":"2019","unstructured":"Zuo LQ, Sun HM, Mao QC et al (2019) Natural scene text recognition based on encoder-decoder framework. IEEE Access 7:62616\u201362623. https:\/\/doi.org\/10.1109\/ACCESS.2019.2916616","journal-title":"IEEE Access"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-023-17671-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-023-17671-1\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-023-17671-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,5,15]],"date-time":"2024-05-15T10:29:31Z","timestamp":1715768971000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-023-17671-1"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,11,30]]},"references-count":94,"journal-issue":{"issue":"18","published-online":{"date-parts":[[2024,5]]}},"alternative-id":["17671"],"URL":"https:\/\/doi.org\/10.1007\/s11042-023-17671-1","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2023,11,30]]},"assertion":[{"value":"25 January 2022","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 September 2023","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"18 November 2023","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 November 2023","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflict of interest. The authors also declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of interest"}}]}}