{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T04:59:21Z","timestamp":1775019561267,"version":"3.50.1"},"reference-count":66,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T00:00:00Z","timestamp":1775001600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T00:00:00Z","timestamp":1775001600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-026-21520-2","type":"journal-article","created":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T03:02:34Z","timestamp":1775012554000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["RT-FCOSH: bridging accuracy and efficiency in low-resolution object detection for autonomous driving"],"prefix":"10.1007","volume":"85","author":[{"ORCID":"https:\/\/orcid.org\/0009-0002-5166-7983","authenticated-orcid":false,"given":"Saad","family":"Mboutayeb","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Aicha","family":"Majda","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Khalid","family":"Zenkouar","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,4,1]]},"reference":[{"key":"21520_CR1","doi-asserted-by":"publisher","first-page":"244","DOI":"10.3311\/PPtr.9464","volume":"44","author":"T Tettamanti","year":"2016","unstructured":"Tettamanti T, Varga I, Szalay Z (2016) Impacts of autonomous cars from a traffic engineering perspective. Period Polytech Transp Eng 44:244\u2013250. https:\/\/doi.org\/10.3311\/PPtr.9464","journal-title":"Period Polytech Transp Eng"},{"key":"21520_CR2","doi-asserted-by":"publisher","first-page":"73","DOI":"10.3390\/photonics6020073","volume":"6","author":"F Sahin","year":"2019","unstructured":"Sahin F (2019) Long-range, high-resolution camera optical design for assisted and autonomous driving. Photonics 6:73. https:\/\/doi.org\/10.3390\/photonics6020073","journal-title":"Photonics"},{"key":"21520_CR3","doi-asserted-by":"publisher","unstructured":"Chavan C, Hembade S, Jadhav G, Komalwad P, Rawat P (2023) Computer vision application analysis based on object detection. Int J Sci Res Eng Manag 07. https:\/\/doi.org\/10.55041\/IJSREM19015","DOI":"10.55041\/IJSREM19015"},{"issue":"3","key":"21520_CR4","doi-asserted-by":"publisher","first-page":"257","DOI":"10.1109\/JPROC.2023.3238524","volume":"111","author":"Z Zou","year":"2023","unstructured":"Zou Z, Chen K, Shi Z, Guo Y, Ye J (2023) Object detection in 20 years: A survey. Proc IEEE 111(3):257\u2013276. https:\/\/doi.org\/10.1109\/JPROC.2023.3238524","journal-title":"Proc IEEE"},{"key":"21520_CR5","doi-asserted-by":"publisher","first-page":"3212","DOI":"10.1109\/TNNLS.2018.2876865","volume":"30","author":"Z-Q Zhao","year":"2019","unstructured":"Zhao Z-Q, Zheng P, Xu S-T, Wu X (2019) Object detection with deep learning: A review. IEEE Trans Neural Netw Learn Syst 30:3212\u20133232. https:\/\/doi.org\/10.1109\/TNNLS.2018.2876865","journal-title":"IEEE Trans Neural Netw Learn Syst"},{"key":"21520_CR6","doi-asserted-by":"publisher","unstructured":"Redmon J, Divvala S, Girshick R, Farhadi A (2015) You only look once: Unified, real-time object detection. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition 2016-Decem, pp. 779\u2013788. https:\/\/doi.org\/10.1109\/CVPR.2016.91","DOI":"10.1109\/CVPR.2016.91"},{"key":"21520_CR7","doi-asserted-by":"publisher","unstructured":"Liu W, Anguelov D, Erhan D, Szegedy C, Reed S, Fu C.-Y, Berg AC (2015) Ssd: Single shot multibox detector. Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics) 9905 LNCS, pp 21\u201337. https:\/\/doi.org\/10.1007\/978-3-319-46448-0_2","DOI":"10.1007\/978-3-319-46448-0_2"},{"key":"21520_CR8","doi-asserted-by":"publisher","unstructured":"Tian Z, Shen C, Chen H, He T (2019) Fcos: Fully convolutional one-stage object detection. In: Proceedings of the IEEE international conference on computer vision 2019-Octob, pp 9626\u20139635. https:\/\/doi.org\/10.1109\/ICCV.2019.00972","DOI":"10.1109\/ICCV.2019.00972"},{"key":"21520_CR9","doi-asserted-by":"publisher","unstructured":"Lin T-Y, Goyal P, Girshick R, He K, Doll\u00e1r P (2017) Focal loss for dense object detection. 13C-NMR of Natural Products, pp 30\u201333. https:\/\/doi.org\/10.1007\/978-1-4615-3288-0_5","DOI":"10.1007\/978-1-4615-3288-0_5"},{"key":"21520_CR10","doi-asserted-by":"publisher","unstructured":"Girshick R, Donahue J, Darrell T, Malik J (2014) Rich feature hierarchies for accurate object detection and semantic segmentation. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition, vol 1, pp 580\u2013587. https:\/\/doi.org\/10.1109\/CVPR.2014.81","DOI":"10.1109\/CVPR.2014.81"},{"key":"21520_CR11","doi-asserted-by":"publisher","unstructured":"Girshick R (2015) Fast r-cnn. In: 2015 IEEE international conference on computer vision (ICCV), vol 2015 Inter, pp 1440\u20131448. IEEE. https:\/\/doi.org\/10.1109\/ICCV.2015.169","DOI":"10.1109\/ICCV.2015.169"},{"key":"21520_CR12","doi-asserted-by":"publisher","first-page":"1137","DOI":"10.1109\/TPAMI.2016.2577031","volume":"39","author":"S Ren","year":"2017","unstructured":"Ren S, He K, Girshick R, Sun J (2017) Faster r-cnn: Towards real-time object detection with region proposal networks. IEEE Trans Pattern Anal Mach Intell 39:1137\u20131149. https:\/\/doi.org\/10.1109\/TPAMI.2016.2577031","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"21520_CR13","doi-asserted-by":"publisher","unstructured":"Cai Z, Vasconcelos N (2018) Cascade r-cnn: Delving into high quality object detection. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition, 6154\u20136162. https:\/\/doi.org\/10.1109\/CVPR.2018.00644","DOI":"10.1109\/CVPR.2018.00644"},{"key":"21520_CR14","doi-asserted-by":"publisher","unstructured":"Chen Q, Wang Y, Yang T, Zhang X, Cheng J, Sun J (2021) You only look one-level feature. In: 2021 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 13034\u201313043. https:\/\/doi.org\/10.1109\/CVPR46437.2021.01284","DOI":"10.1109\/CVPR46437.2021.01284"},{"key":"21520_CR15","doi-asserted-by":"publisher","first-page":"200324","DOI":"10.1016\/j.iswa.2023.200324","volume":"21","author":"S Mboutayeb","year":"2024","unstructured":"Mboutayeb S, Majda A, Zenkouar K, Nikolov NS (2024) Fcosh: A novel single-head fcos for faster object detection in autonomous-driving systems. Intell Syst Appl 21:200324. https:\/\/doi.org\/10.1016\/j.iswa.2023.200324","journal-title":"Intell Syst Appl"},{"key":"21520_CR16","doi-asserted-by":"crossref","unstructured":"Yang C, Huang Z, Wang N (2021) Querydet: Cascaded sparse query for accelerating high-resolution small object detection","DOI":"10.1109\/CVPR52688.2022.01330"},{"key":"21520_CR17","doi-asserted-by":"publisher","unstructured":"Yu F, Chen H, Wang X, Xian W, Chen Y, Liu F, Madhavan V, Darrell T (2018) Bdd100k: A diverse driving dataset for heterogeneous multitask learning. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition, pp 2633\u20132642. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00271","DOI":"10.1109\/CVPR42600.2020.00271"},{"key":"21520_CR18","doi-asserted-by":"publisher","first-page":"207","DOI":"10.1109\/TIP.2020.3034487","volume":"30","author":"Y Pang","year":"2021","unstructured":"Pang Y, Cao J, Li Y, Xie J, Sun H, Gong J (2021) Tju-dhd: A diverse high-resolution dataset for object detection. IEEE Trans Image Process 30:207\u2013219. https:\/\/doi.org\/10.1109\/TIP.2020.3034487","journal-title":"IEEE Trans Image Process"},{"issue":"6","key":"21520_CR19","doi-asserted-by":"publisher","first-page":"1856","DOI":"10.1109\/TMI.2019.2959609","volume":"39","author":"Z Zhou","year":"2020","unstructured":"Zhou Z, Siddiquee MMR, Tajbakhsh N, Liang J (2020) Unet++: Redesigning skip connections to exploit multiscale features in image segmentation. IEEE Trans Med Imaging 39(6):1856\u20131867. https:\/\/doi.org\/10.1109\/TMI.2019.2959609","journal-title":"IEEE Trans Med Imaging"},{"key":"21520_CR20","doi-asserted-by":"publisher","unstructured":"Liu Z, Lin Y, Cao Y, Hu H, Wei Y, Zhang Z, Lin S, Guo B (2021) Swin transformer: Hierarchical vision transformer using shifted windows. In: 2021 IEEE\/CVF international conference on computer vision (ICCV), pp 9992\u201310002. https:\/\/doi.org\/10.1109\/ICCV48922.2021.00986","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"21520_CR21","unstructured":"Redmon J, Farhadi A (2018) Yolov3: An incremental improvement"},{"key":"21520_CR22","doi-asserted-by":"crossref","unstructured":"Lin T-Y, Doll\u00e1r P, Girshick R, He K, Hariharan B, Belongie S (2016) Feature pyramid networks for object detection","DOI":"10.1109\/CVPR.2017.106"},{"key":"21520_CR23","doi-asserted-by":"publisher","first-page":"7389","DOI":"10.1109\/TIP.2020.3002345","volume":"29","author":"T Kong","year":"2020","unstructured":"Kong T, Sun F, Liu H, Jiang Y, Li L, Shi J (2020) Foveabox: Beyound anchor-based object detection. IEEE Trans Image Process 29:7389\u20137398. https:\/\/doi.org\/10.1109\/TIP.2020.3002345","journal-title":"IEEE Trans Image Process"},{"key":"21520_CR24","doi-asserted-by":"publisher","unstructured":"Wang N, Gao Y, Chen H, Wang P, Tian Z, Shen C, Zhang Y (2020) Nas-fcos: Fast neural architecture search for object detection. In: 2020 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 11940\u201311948. https:\/\/doi.org\/10.1109\/CVPR42600.2020.01196","DOI":"10.1109\/CVPR42600.2020.01196"},{"key":"21520_CR25","doi-asserted-by":"publisher","first-page":"83535","DOI":"10.1007\/s11042-024-18872-y","volume":"83","author":"A Vijayakumar","year":"2024","unstructured":"Vijayakumar A, Vairavasundaram S (2024) Yolo-based object detection models: A review and its applications. Multimed Tools Appl 83:83535\u201383574","journal-title":"Multimed Tools Appl"},{"key":"21520_CR26","doi-asserted-by":"publisher","unstructured":"Jocher G, Chaurasia A, Stoken A, Borovec J, NanoCode012, Kwon Y, TaoXie, Michael K, Fang J, imyhxy, Lorna, Wong C, Yifu Z, V A, Montes D, Wang Z, Fati C, Nadar J, Laughing, UnglvKitDe, tkianai, yxNONG, Skalski P, Hogan A, Strobel M, Jain M, Mammana L. Xylieong: ultralytics\/yolov5: V6.2 - YOLOv5 Classification Models, Apple M1, Reproducibility, ClearML and Deci.ai Integrations. https:\/\/doi.org\/10.5281\/zenodo.7002879","DOI":"10.5281\/zenodo.7002879"},{"key":"21520_CR27","unstructured":"Ge Z, Liu S, Wang F, Li Z, Sun J (2021) Yolox: Exceeding yolo series in 2021"},{"key":"21520_CR28","unstructured":"Li C, Li L, Jiang H, Weng K, Geng Y, Li L, Ke Z, Li Q, Cheng M, Nie W, Li Y, Zhang B, Liang Y, Zhou L, Xu X, Chu X, Wei X, Wei X (2022) YOLOv6: A single-stage object detection framework for industrial applications. arxiv:2209.02976"},{"key":"21520_CR29","doi-asserted-by":"publisher","unstructured":"Wang C-Y, Bochkovskiy A, Liao H-YM (2023) Yolov7: Trainable bag-of-freebies sets new state-of-the-art for real-time object detectors. In: 2023 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 7464\u20137475. https:\/\/doi.org\/10.1109\/CVPR52729.2023.00721","DOI":"10.1109\/CVPR52729.2023.00721"},{"key":"21520_CR30","unstructured":"Jocher G, Chaurasia A, Qiu J. Ultralytics YOLOv8. https:\/\/github.com\/ultralytics\/ultralytics"},{"key":"21520_CR31","doi-asserted-by":"crossref","unstructured":"Wang C-Y, Liao H-YM (2024) Yolov9: Learning what you want to learn using programmable gradient information","DOI":"10.1007\/978-3-031-72751-1_1"},{"key":"21520_CR32","doi-asserted-by":"crossref","unstructured":"Zhang S, Chi C, Yao Y, Lei Z, Li SZ (2019) Bridging the gap between anchor-based and anchor-free detection via adaptive training sample selection","DOI":"10.1109\/CVPR42600.2020.00978"},{"key":"21520_CR33","doi-asserted-by":"publisher","first-page":"642","DOI":"10.1007\/s11263-019-01204-1","volume":"128","author":"H Law","year":"2018","unstructured":"Law H, Deng J (2018) Cornernet: Detecting objects as paired keypoints. Int J Comput Vision 128:642\u2013656. https:\/\/doi.org\/10.1007\/s11263-019-01204-1","journal-title":"Int J Comput Vision"},{"key":"21520_CR34","doi-asserted-by":"publisher","unstructured":"Duan K, Bai S, Xie L, Qi H, Huang Q, Tian Q (2019) Centernet: Keypoint triplets for object detection. In: Proceedings of the IEEE international conference on computer vision 2019-Octob, pp 6568\u20136577. https:\/\/doi.org\/10.1109\/ICCV.2019.00667","DOI":"10.1109\/ICCV.2019.00667"},{"key":"21520_CR35","doi-asserted-by":"publisher","unstructured":"Zhou X, Zhuo J, Kr\u00e4henb\u00fchl P (2019) Bottom-up object detection by grouping extreme and center points. In: Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition 2019-June, pp 850\u2013859. https:\/\/doi.org\/10.1109\/CVPR.2019.00094","DOI":"10.1109\/CVPR.2019.00094"},{"key":"21520_CR36","doi-asserted-by":"publisher","unstructured":"Newell A, Yang K, Deng J (2016) Stacked hourglass networks for human pose estimation. Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics) 9912 LNCS, pp 483\u2013499. https:\/\/doi.org\/10.1007\/978-3-319-46484-8_29","DOI":"10.1007\/978-3-319-46484-8_29"},{"issue":"25","key":"21520_CR37","doi-asserted-by":"publisher","first-page":"30415","DOI":"10.1007\/s11042-024-20402-9","volume":"84","author":"X Huakai","year":"2025","unstructured":"Huakai X (2025) Feature fusion means a lot to detrs. Multimed Tools Appl 84(25):30415\u201330429","journal-title":"Multimed Tools Appl"},{"key":"21520_CR38","doi-asserted-by":"publisher","unstructured":"Carion N, Massa F, Synnaeve G, Usunier N, Kirillov A, Zagoruyko S (2020) End-to-end object detection with transformers. In: Computer Vision \u2013 ECCV 2020: 16th European Conference, Glasgow, UK, August 23\u201328, 2020, Proceedings, Part I, pp 213\u2013229. Springer, Berlin, Heidelberg.https:\/\/doi.org\/10.1007\/978-3-030-58452-8_13","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"21520_CR39","unstructured":"Zhu X, Su W, Lu L, Li B, Wang X, Dai J (2020) Deformable detr: Deformable transformers for end-to-end object detection. arXiv:2010.04159"},{"key":"21520_CR40","doi-asserted-by":"publisher","unstructured":"Meng D, Chen X, Fan Z, Zeng G, Li H, Yuan Y, Sun L, Wang J (2021) Conditional detr for fast training convergence. In: 2021 IEEE\/CVF international conference on computer vision (ICCV), pp 3631\u20133640. https:\/\/doi.org\/10.1109\/ICCV48922.2021.00363","DOI":"10.1109\/ICCV48922.2021.00363"},{"key":"21520_CR41","unstructured":"Liu S, Li F, Zhang H, Yang X.B, Qi X, Su H, Zhu J, Zhang L (2022) Dab-detr: Dynamic anchor boxes are better queries for detr. ArXiv:2201.12329"},{"key":"21520_CR42","unstructured":"Zhang H, Li F, Liu S, Zhang L, Su H, Zhu J, Ni LM, Shum H-Y (2022) Dino: Detr with improved denoising anchor boxes for end-to-end object detection. arXiv preprint arXiv:2203.03605"},{"issue":"11","key":"21520_CR43","doi-asserted-by":"publisher","first-page":"12994","DOI":"10.1109\/TII.2024.3431044","volume":"20","author":"S Cheng","year":"2024","unstructured":"Cheng S, Song J, Zhou M, Wei X, Pu H, Luo J, Jia W (2024) Ef-detr: A lightweight transformer-based object detector with an encoder-free neck. IEEE Trans Industr Inf 20(11):12994\u201313002. https:\/\/doi.org\/10.1109\/TII.2024.3431044","journal-title":"IEEE Trans Industr Inf"},{"issue":"36","key":"21520_CR44","doi-asserted-by":"publisher","first-page":"84231","DOI":"10.1007\/s11042-024-19087-x","volume":"83","author":"AK Tiwari","year":"2024","unstructured":"Tiwari AK (2024) Pattanaik M, Sharma G: Low-light detection transformer (ldetr): object detection in low-light and adverse weather conditions. Multimed Tools Appl 83(36):84231\u201384248","journal-title":"Multimed Tools Appl"},{"issue":"41","key":"21520_CR45","doi-asserted-by":"publisher","first-page":"88645","DOI":"10.1007\/s11042-024-18866-w","volume":"83","author":"GKJ Iqra","year":"2024","unstructured":"Iqra GKJ, Javed M (2024) Small object detection in diverse application landscapes: a survey. Multimed Tools Appl 83(41):88645\u201388680","journal-title":"Multimed Tools Appl"},{"issue":"28","key":"21520_CR46","doi-asserted-by":"publisher","first-page":"70727","DOI":"10.1007\/s11042-023-18100-z","volume":"83","author":"L Cui-jin","year":"2024","unstructured":"Cui-jin L, Zhong Q, Sheng-ye W (2024) Ghafnet: Global-context hierarchical attention fusion method for traffic object detection. Multimed Tools Appl 83(28):70727\u201370748","journal-title":"Multimed Tools Appl"},{"key":"21520_CR47","doi-asserted-by":"crossref","unstructured":"Xu S, Zheng S, Xu W, Xu R, Wang C, Zhang J, Teng X, Li A, Guo L (2024) Hcf-net: Hierarchical context fusion network for infrared small object detection. In: 2024 IEEE international conference on multimedia and expo (ICME). IEEE, pp 1\u20136","DOI":"10.1109\/ICME57554.2024.10687431"},{"key":"21520_CR48","doi-asserted-by":"publisher","DOI":"10.1016\/j.jvcir.2023.103752","volume":"90","author":"M Wang","year":"2023","unstructured":"Wang M, Yang W, Wang L, Chen D, Wei F, KeZiErBieKe H, Liao Y (2023) Fe-yolov5: Feature enhancement network based on yolov5 for small object detection. J Vis Commun Image Represent 90:103752. https:\/\/doi.org\/10.1016\/j.jvcir.2023.103752","journal-title":"J Vis Commun Image Represent"},{"issue":"22","key":"21520_CR49","doi-asserted-by":"publisher","first-page":"61239","DOI":"10.1007\/s11042-023-17818-0","volume":"83","author":"F Yang","year":"2024","unstructured":"Yang F, Zhou J, Chen Y, Liao J, Yang M (2024) Msf-yolo: A multi-scale features fusion-based method for small object detection. Multimed Tools Appl 83(22):61239\u201361260","journal-title":"Multimed Tools Appl"},{"key":"21520_CR50","doi-asserted-by":"crossref","unstructured":"Sunkara R, Luo T (2022) No more strided convolutions or pooling: A new cnn building block for low-resolution images and small objects. In: Joint European conference on machine learning and knowledge discovery in databases. Springer, pp 443\u2013459","DOI":"10.1007\/978-3-031-26409-2_27"},{"key":"21520_CR51","doi-asserted-by":"crossref","unstructured":"Tang S, Zhang S, Fang Y (2024) Hic-yolov5: Improved yolov5 for small object detection. In: 2024 IEEE international conference on robotics and automation (ICRA). IEEE, pp 6614\u20136619","DOI":"10.1109\/ICRA57147.2024.10610273"},{"key":"21520_CR52","doi-asserted-by":"crossref","unstructured":"Li D, Hu J, Wang C, Li X, She Q, Zhu L, Zhang T, Chen Q (2021) Involution: Inverting the inherence of convolution for visual recognition. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 12321\u201312330","DOI":"10.1109\/CVPR46437.2021.01214"},{"key":"21520_CR53","unstructured":"Rekavandi A.M, Rashidi S, Boussaid F, Hoefs S, Akbas E, et al (2023) Transformers in small object detection: A benchmark and survey of state-of-the-art. arXiv preprint arXiv:2309.04902"},{"key":"21520_CR54","doi-asserted-by":"crossref","unstructured":"Lin C-L, Chen Y-L, Lin Y-C (2022) Small objects detection using transformer-cnn. Available at SSRN 4232850","DOI":"10.2139\/ssrn.4232850"},{"key":"21520_CR55","doi-asserted-by":"publisher","unstructured":"Dubois E, Mitiche A (1984) Review of \u2019digital picture processing,\u2019 2nd edn. (rosenfeld, a., and kak, a.c.; 1982). IEEE Trans Inf Theory 30:694\u2013695. https:\/\/doi.org\/10.1109\/TIT.1984.1056929","DOI":"10.1109\/TIT.1984.1056929"},{"key":"21520_CR56","doi-asserted-by":"publisher","first-page":"1153","DOI":"10.1109\/TASSP.1981.1163711","volume":"29","author":"R Keys","year":"1981","unstructured":"Keys R (1981) Cubic convolution interpolation for digital image processing. IEEE Trans Acoust Speech Signal Process 29:1153\u20131160. https:\/\/doi.org\/10.1109\/TASSP.1981.1163711","journal-title":"IEEE Trans Acoust Speech Signal Process"},{"key":"21520_CR57","doi-asserted-by":"publisher","unstructured":"Zeiler M.D, Krishnan D, Taylor G.W, Fergus R (2010) Deconvolutional networks. In: 2010 IEEE computer society conference on computer vision and pattern recognition, pp 2528\u20132535. IEEE. https:\/\/doi.org\/10.1109\/CVPR.2010.5539957","DOI":"10.1109\/CVPR.2010.5539957"},{"key":"21520_CR58","doi-asserted-by":"publisher","unstructured":"Shi W, Caballero J, Huszar F, Totz J, Aitken A.P, Bishop R, Rueckert D, Wang Z (2016) Real-time single image and video super-resolution using an efficient sub-pixel convolutional neural network. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition 2016-December, pp 1874\u20131883. https:\/\/doi.org\/10.1109\/CVPR.2016.207","DOI":"10.1109\/CVPR.2016.207"},{"key":"21520_CR59","doi-asserted-by":"publisher","unstructured":"Wang J, Chen K, Xu R, Liu Z, Loy C.C, Lin D (2019) Carafe: Content-aware reassembly of features. In: IEEE International Conference on Computer Vision 2019-October, pp 3007\u20133016. https:\/\/doi.org\/10.1109\/ICCV.2019.00310","DOI":"10.1109\/ICCV.2019.00310"},{"key":"21520_CR60","doi-asserted-by":"publisher","unstructured":"Lin T-Y, Maire M, Belongie S, Hays J, Perona P, Ramanan D, Doll\u00e1r P, Zitnick CL (2014) Microsoft COCO: Common Objects in Context, vol. 8693 LNCS, pp 740\u2013755. https:\/\/doi.org\/10.1007\/978-3-319-10602-1_48","DOI":"10.1007\/978-3-319-10602-1_48"},{"key":"21520_CR61","unstructured":"Ruder S (2016) An overview of gradient descent optimization algorithms. ArXiv:1609.04747"},{"key":"21520_CR62","unstructured":"Loshchilov I, Hutter F (2019) Decoupled weight decay regularization. arxiv:1711.05101"},{"key":"21520_CR63","unstructured":"Chen K, Wang J, Pang J, Cao Y, Xiong Y, Li X, Sun S, Feng W, Liu Z, Xu J, Zhang Z, Cheng D, Zhu C, Cheng T, Zhao Q, Li B, Lu X, Zhu R, Wu Y, Dai J, Wang J, Shi J, Ouyang W, Loy CC, Lin D (2019) Mmdetection: Open mmlab detection toolbox and benchmark"},{"key":"21520_CR64","doi-asserted-by":"publisher","unstructured":"Bodla N, Singh B, Chellappa R, Davis LS (2017) Soft-nms \u2013 improving object detection with one line of code. In: Proceedings of the IEEE international conference on computer vision 2017-Octob, pp 5562\u20135570. https:\/\/doi.org\/10.1109\/ICCV.2017.593","DOI":"10.1109\/ICCV.2017.593"},{"key":"21520_CR65","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3508391","volume":"21","author":"E Jeong","year":"2022","unstructured":"Jeong E, Kim J, Ha S (2022) Tensorrt-based framework and optimization methodology for deep learning inference on jetson boards. ACM Trans Embedded Comput Syst 21:1\u201326. https:\/\/doi.org\/10.1145\/3508391","journal-title":"ACM Trans Embedded Comput Syst"},{"key":"21520_CR66","unstructured":"Seo D, Yang H, Park S, Kim H (2024) Elastic-detr: Making image resolution learnable with content-specific network prediction. ArXiv:2412.06341"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-026-21520-2.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-026-21520-2","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-026-21520-2.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T03:02:58Z","timestamp":1775012578000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-026-21520-2"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,4,1]]},"references-count":66,"journal-issue":{"issue":"4","published-online":{"date-parts":[[2026,4]]}},"alternative-id":["21520"],"URL":"https:\/\/doi.org\/10.1007\/s11042-026-21520-2","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,4,1]]},"assertion":[{"value":"8 July 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"12 February 2026","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"23 March 2026","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"1 April 2026","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"Not applicable.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical Approval"}},{"value":"Not applicable.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Consent to Participate"}},{"value":"Not applicable.","order":4,"name":"Ethics","group":{"name":"EthicsHeading","label":"Consent to Publish"}},{"value":"The authors declare that they have no competing interests.","order":5,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing Interests"}}],"article-number":"325"}}