{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,31]],"date-time":"2026-03-31T06:28:46Z","timestamp":1774938526374,"version":"3.50.1"},"reference-count":65,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2026,2,1]],"date-time":"2026-02-01T00:00:00Z","timestamp":1769904000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,2,1]],"date-time":"2026-02-01T00:00:00Z","timestamp":1769904000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100002858","name":"China Postdoctoral Science Foundation","doi-asserted-by":"publisher","award":["Grant No. 259822"],"award-info":[{"award-number":["Grant No. 259822"]}],"id":[{"id":"10.13039\/501100002858","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100012152","name":"National Postdoctoral Program for Innovative Talents","doi-asserted-by":"publisher","award":["Grant No. BX2020010"],"award-info":[{"award-number":["Grant No. BX2020010"]}],"id":[{"id":"10.13039\/501100012152","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/100014718","name":"Innovative Research Group Project of the National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["Grant No. 62206077"],"award-info":[{"award-number":["Grant No. 62206077"]}],"id":[{"id":"10.13039\/100014718","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/100014718","name":"Innovative Research Group Project of the National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["No. 61976070"],"award-info":[{"award-number":["No. 61976070"]}],"id":[{"id":"10.13039\/100014718","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100005046","name":"Natural Science Foundation of Heilongjiang Province","doi-asserted-by":"publisher","award":["LH2021F024"],"award-info":[{"award-number":["LH2021F024"]}],"id":[{"id":"10.13039\/501100005046","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2026,2]]},"DOI":"10.1007\/s10489-026-07139-8","type":"journal-article","created":{"date-parts":[[2026,2,14]],"date-time":"2026-02-14T06:58:44Z","timestamp":1771052324000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Stream-DINO: exploring DETR-based online object detection with streaming perception"],"prefix":"10.1007","volume":"56","author":[{"given":"Man","family":"Zhang","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0437-7337","authenticated-orcid":false,"given":"Yongqiang","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Rui","family":"Tian","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yin","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zian","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jinwei","family":"Sun","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,2,14]]},"reference":[{"key":"7139_CR1","doi-asserted-by":"publisher","first-page":"58443","DOI":"10.1109\/ACCESS.2020.2983149","volume":"8","author":"E Yurtsever","year":"2020","unstructured":"Yurtsever E, Lambert J, Carballo A, Takeda K (2020) A survey of autonomous driving: Common practices and emerging technologies. IEEE access 8:58443\u201358469","journal-title":"IEEE access"},{"key":"7139_CR2","doi-asserted-by":"publisher","first-page":"100057","DOI":"10.1016\/j.array.2021.100057","volume":"10","author":"A Gupta","year":"2021","unstructured":"Gupta A, Anpalagan A, Guan L, Khwaja AS (2021) Deep learning for object detection and scene perception in self-driving cars: Survey, challenges, and open issues. Array 10:100057","journal-title":"Array"},{"issue":"3","key":"7139_CR3","doi-asserted-by":"publisher","first-page":"1403","DOI":"10.1109\/TCSVT.2021.3072207","volume":"32","author":"T Zhang","year":"2021","unstructured":"Zhang T, Liu X, Zhang Q, Han J (2021) Siamcda: Complementarity-and distractor-aware rgb-t tracking based on siamese network. IEEE Trans Circuits Syst Video Technol 32(3):1403\u20131417","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"7139_CR4","doi-asserted-by":"crossref","unstructured":"Liu W, Li W, Zhu J, Cui M, Xie X, Zhang L (2023) Improving nighttime driving-scene segmentation via dual image-adaptive learnable filters. IEEE Trans Circ Syst Video Technol","DOI":"10.1109\/TCSVT.2023.3260240"},{"issue":"6","key":"7139_CR5","doi-asserted-by":"publisher","first-page":"4689","DOI":"10.1007\/s10489-024-05403-3","volume":"54","author":"Y Zhang","year":"2024","unstructured":"Zhang Y, Tian R, Zhang Y, Zhang Z, Bai Y, Ding M, Zuo W (2024) R-ccf: region-aware continual contrastive fusion for weakly supervised object detection. Appl Intell 54(6):4689\u20134712","journal-title":"Appl Intell"},{"key":"7139_CR6","unstructured":"Ren S, He K, Girshick R, Sun J (2015) Faster r-cnn: Towards real-time object detection with region proposal networks. Adv Neural Inf Process Syst 28"},{"key":"7139_CR7","doi-asserted-by":"crossref","unstructured":"He K, Gkioxari G, Doll\u00e1r P, Girshick R (2017) Mask r-cnn. In: Proceedings of the IEEE international conference on computer vision, pp 2961\u20132969","DOI":"10.1109\/ICCV.2017.322"},{"key":"7139_CR8","doi-asserted-by":"crossref","unstructured":"Liu W, Anguelov D, Erhan D, Szegedy C, Reed S, Fu CY, Berg AC (2016) Ssd: Single shot multibox detector. In: Computer Vision-ECCV 2016: 14th European Conference, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part I 14, Springer, pp 21\u201337","DOI":"10.1007\/978-3-319-46448-0_2"},{"key":"7139_CR9","doi-asserted-by":"crossref","unstructured":"Cai Z, Vasconcelos N (2018) Cascade r-cnn: Delving into high quality object detection. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 6154\u20136162","DOI":"10.1109\/CVPR.2018.00644"},{"key":"7139_CR10","doi-asserted-by":"crossref","unstructured":"Gavrilescu R, Zet C, Fo\u0219al\u0103u C, Skoczylas M, Cotovanu D (2018) Faster r-cnn: an approach to real-time object detection. In: 2018 International conference and exposition on electrical and power engineering (EPE), IEEE, pp 0165\u20130168","DOI":"10.1109\/ICEPE.2018.8559776"},{"issue":"6","key":"7139_CR11","doi-asserted-by":"publisher","first-page":"9243","DOI":"10.1007\/s11042-022-13644-y","volume":"82","author":"T Diwan","year":"2023","unstructured":"Diwan T, Anirudh G, Tembhurne JV (2023) Object detection using yolo: Challenges, architectural successors, datasets and applications. Multimed Tools Appl 82(6):9243\u20139275","journal-title":"Multimed Tools Appl"},{"key":"7139_CR12","unstructured":"Ge Z, Liu S, Wang F, Li Z, Sun J (2021) Yolox: Exceeding yolo series in 2021. arXiv:2107.08430"},{"key":"7139_CR13","doi-asserted-by":"crossref","unstructured":"Wang CY, Bochkovskiy A, Liao HYM (2023) Yolov7: Trainable bag-of-freebies sets new state-of-the-art for real-time object detectors. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 7464\u20137475","DOI":"10.1109\/CVPR52729.2023.00721"},{"key":"7139_CR14","doi-asserted-by":"crossref","unstructured":"Redmon J, Divvala S, Girshick R, Farhadi A (2016) You only look once: Unified, real-time object detection. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 779\u2013788","DOI":"10.1109\/CVPR.2016.91"},{"key":"7139_CR15","doi-asserted-by":"crossref","unstructured":"Sultana F, Sufian A, Dutta P (2020) A review of object detection models based on convolutional neural network. Intell Comput: image Process Based Appl 1\u201316","DOI":"10.1007\/978-981-15-4288-6_1"},{"key":"7139_CR16","doi-asserted-by":"crossref","unstructured":"Liu S, Huang D, Wang Y (2019) Adaptive nms: Refining pedestrian detection in a crowd. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 6459\u20136468","DOI":"10.1109\/CVPR.2019.00662"},{"key":"7139_CR17","doi-asserted-by":"crossref","unstructured":"Zhu Y, Zhang Y, Ding M, Zuo W (2022) Uncertainty-aware graph-guided weakly supervised object detection. IEEE Trans Circ Syst Video Technol","DOI":"10.1109\/TCSVT.2022.3232487"},{"key":"7139_CR18","doi-asserted-by":"crossref","unstructured":"Li M, Wang YX, Ramanan D (2020) Towards streaming perception. In: 16th European Conference on Computer Vision, ECCV 2020, Springer, pp 473\u2013488","DOI":"10.1007\/978-3-030-58536-5_28"},{"key":"7139_CR19","doi-asserted-by":"crossref","unstructured":"Dong N, Zhang Y, Ding M, Lee GH (2024) Towards non co-occurrence incremental object detection with unlabeled in-the-wild data. Intern J Comput Vis 1\u201318","DOI":"10.1007\/s11263-024-02048-0"},{"key":"7139_CR20","doi-asserted-by":"crossref","unstructured":"Zhang Y, Zhang Y, Zhang Z, Zhang M, Tian R, Ding M (2024a) Isp-teacher: Image signal process with disentanglement regularization for unsupervised domain adaptive dark object detection. In: Proceedings of the AAAI conference on artificial intelligence, vol 38, pp 7387\u20137395","DOI":"10.1609\/aaai.v38i7.28569"},{"issue":"6","key":"7139_CR21","doi-asserted-by":"publisher","first-page":"4689","DOI":"10.1007\/s10489-024-05403-3","volume":"54","author":"Y Zhang","year":"2024","unstructured":"Zhang Y, Tian R, Zhang Y, Zhang Z, Bai Y, Ding M, Zuo W (2024) R-ccf: region-aware continual contrastive fusion for weakly supervised object detection. Appl Intell 54(6):4689\u20134712","journal-title":"Appl Intell"},{"key":"7139_CR22","doi-asserted-by":"publisher","first-page":"110111","DOI":"10.1016\/j.patcog.2023.110111","volume":"147","author":"Z Zhang","year":"2024","unstructured":"Zhang Z, Zhang Y, Zhang Y, Tian R, Ding M (2024) Vital information is only worth one thumbnail: Towards efficient human pose estimation. Pattern Recogn 147:110111","journal-title":"Pattern Recogn"},{"key":"7139_CR23","doi-asserted-by":"crossref","unstructured":"Dong N, Zhang Y, Ding M, Lee GH (2023) Boosting long-tailed object detection via step-wise learning on smooth-tail data. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 6940\u20136949","DOI":"10.1109\/ICCV51070.2023.00639"},{"key":"7139_CR24","unstructured":"Zhu X, Su W, Lu L, Li B, Wang X, Dai J (2020) Deformable detr: Deformable transformers for end-to-end object detection. In: International conference on learning representations"},{"key":"7139_CR25","doi-asserted-by":"crossref","unstructured":"Carion N, Massa F, Synnaeve G, Usunier N, Kirillov A, Zagoruyko S (2020) End-to-end object detection with transformers. In: European conference on computer vision, Springer, pp 213\u2013229","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"7139_CR26","doi-asserted-by":"crossref","unstructured":"Dai X, Chen Y, Yang J, Zhang P, Yuan L, Zhang L (2021) Dynamic detr: End-to-end object detection with dynamic attention. In: 2021 IEEE\/CVF International conference on computer vision (ICCV), IEEE Computer Society, pp 2968\u20132977","DOI":"10.1109\/ICCV48922.2021.00298"},{"key":"7139_CR27","doi-asserted-by":"crossref","unstructured":"Li F, Zhang H, Liu S, Guo J, Ni LM, Zhang L (2022) Dn-detr: Accelerate detr training by introducing query denoising. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 13619\u201313627","DOI":"10.1109\/CVPR52688.2022.01325"},{"key":"7139_CR28","unstructured":"Zhang H, Li F, Liu S, Zhang L, Su H, Zhu J, Ni LM, Shum HY (2022) Dino: Detr with improved denoising anchor boxes for end-to-end object detection. arXiv preprint arXiv:2203.03605"},{"key":"7139_CR29","unstructured":"Liu S, Li F, Zhang H, Yang X, Qi X, Su H, Zhu J, Zhang L (2022) Dab-detr: Dynamic anchor boxes are better queries for detr. arXiv:2201.12329"},{"key":"7139_CR30","unstructured":"Long X, Deng K, Wang G, Zhang Y, Dang Q, Gao Y, Shen H, Ren J, Han S, Ding E et al (2020) Pp-yolo: An effective and efficient implementation of object detector. arXiv:2007.12099"},{"issue":"4","key":"7139_CR31","doi-asserted-by":"publisher","first-page":"983","DOI":"10.1109\/TCSVT.2019.2898559","volume":"30","author":"Y Zhang","year":"2019","unstructured":"Zhang Y, Ding M, Bai Y, Xu M, Ghanem B (2019) Beyond weakly supervised: Pseudo ground truths mining for missing bounding-boxes object detection. IEEE Trans Circuits Syst Video Technol 30(4):983\u2013997","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"7139_CR32","doi-asserted-by":"crossref","unstructured":"Yang J, Liu S, Li Z, Li X, Sun J (2022) Real-time object detection for streaming perception. In: 2022 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), IEEE Computer Society, pp 5375\u20135385","DOI":"10.1109\/CVPR52688.2022.00531"},{"key":"7139_CR33","doi-asserted-by":"crossref","unstructured":"Zhao Y, Lv W, Xu S, Wei J, Wang G, Dang Q, Liu Y, Chen J (2023) Detrs beat yolos on real-time object detection. arXiv:2304.08069","DOI":"10.1109\/CVPR52733.2024.01605"},{"key":"7139_CR34","doi-asserted-by":"crossref","unstructured":"Liu S, Qi L, Qin H, Shi J, Jia J (2018) Path aggregation network for instance segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 8759\u20138768","DOI":"10.1109\/CVPR.2018.00913"},{"key":"7139_CR35","doi-asserted-by":"crossref","unstructured":"Lin TY, Goyal P, Girshick R, He K, Doll\u00e1r P (2017) Focal loss for dense object detection. In: Proceedings of the IEEE international conference on computer vision, pp 2980\u20132988","DOI":"10.1109\/ICCV.2017.324"},{"issue":"3","key":"7139_CR36","doi-asserted-by":"publisher","first-page":"2923","DOI":"10.1007\/s10489-022-03686-y","volume":"53","author":"Z Mao","year":"2023","unstructured":"Mao Z, Zhou Y, Sun J, Wu H, Pan F, Ahmad B (2023) Weakly-supervised object localization with gradient-pyramid feature. Appl Intell 53(3):2923\u20132935","journal-title":"Appl Intell"},{"issue":"5","key":"7139_CR37","doi-asserted-by":"publisher","first-page":"2251","DOI":"10.1109\/TNNLS.2020.3004819","volume":"32","author":"Y Zhang","year":"2020","unstructured":"Zhang Y, Bai Y, Ding M, Xu S, Ghanem B (2020) Kgsnet: key-point-guided super-resolution network for pedestrian detection in the wild. IEEE Trans Neural Netw Learn Syst 32(5):2251\u20132265","journal-title":"IEEE Trans Neural Netw Learn Syst"},{"issue":"6","key":"7139_CR38","doi-asserted-by":"publisher","first-page":"1810","DOI":"10.1007\/s11263-020-01301-6","volume":"128","author":"Y Zhang","year":"2020","unstructured":"Zhang Y, Bai Y, Ding M, Ghanem B (2020) Multi-task generative adversarial network for detecting small objects in the wild. Int J Comput Vision 128(6):1810\u20131828","journal-title":"Int J Comput Vision"},{"key":"7139_CR39","doi-asserted-by":"publisher","first-page":"109424","DOI":"10.1016\/j.patcog.2023.109424","volume":"138","author":"Y Zhang","year":"2023","unstructured":"Zhang Y, Zhang Y, Tian R, Zhang Z, Bai Y, Zuo W, Ding M (2023) Thumbdet: One thumbnail image is enough for object detection. Pattern Recogn 138:109424","journal-title":"Pattern Recogn"},{"key":"7139_CR40","doi-asserted-by":"crossref","unstructured":"Neubeck A, Van Gool L (2006) Efficient non-maximum suppression. In: 18th international conference on pattern recognition (ICPR\u201906), IEEE, vol 3, pp 850\u2013855","DOI":"10.1109\/ICPR.2006.479"},{"issue":"3","key":"7139_CR41","doi-asserted-by":"publisher","first-page":"1320","DOI":"10.1109\/TCSVT.2022.3210207","volume":"33","author":"J Leng","year":"2022","unstructured":"Leng J, Mo M, Zhou Y, Gao C, Li W, Gao X (2022) Pareto refocusing for drone-view object detection. IEEE Trans Circuits Syst Video Technol 33(3):1320\u20131334","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"issue":"10","key":"7139_CR42","doi-asserted-by":"publisher","first-page":"105023","DOI":"10.1088\/1361-6501\/ad633d","volume":"35","author":"H Tao","year":"2024","unstructured":"Tao H, Zheng Y, Wang Y, Qiu J, Stojanovic V (2024) Enhanced feature extraction yolo industrial small object detection algorithm based on receptive-field attention and multi-scale features. Meas Sci Technol 35(10):105023","journal-title":"Meas Sci Technol"},{"issue":"3","key":"7139_CR43","doi-asserted-by":"publisher","first-page":"035004","DOI":"10.1088\/1361-6501\/adb2ad","volume":"36","author":"H Tao","year":"2025","unstructured":"Tao H, Huang Z, Wang Y, Qiu J, Vladimir S (2025) Efficient feature fusion network for small objects detection of traffic signs based on cross-dimensional and dual-domain information. Meas Sci Technol 36(3):035004","journal-title":"Meas Sci Technol"},{"key":"7139_CR44","doi-asserted-by":"publisher","first-page":"109402","DOI":"10.1016\/j.engappai.2024.109402","volume":"138","author":"Y Sun","year":"2024","unstructured":"Sun Y, Tao H, Stojanovic V (2024) Autoregressive data generation method based on wavelet packet transform and cascaded stochastic quantization for bearing fault diagnosis under unbalanced samples. Eng Appl Artif Intell 138:109402","journal-title":"Eng Appl Artif Intell"},{"key":"7139_CR45","doi-asserted-by":"crossref","unstructured":"Abed A, Akrout B, Amous I (2022) Shoppers interaction classification based on an improved densenet model using rgb-d data. In: 2022 8th International conference on systems and informatics (ICSAI), IEEE, pp 1\u20136","DOI":"10.1109\/ICSAI57119.2022.10005508"},{"issue":"31","key":"7139_CR46","doi-asserted-by":"publisher","first-page":"19365","DOI":"10.1007\/s00521-024-10239-6","volume":"36","author":"A Abed","year":"2024","unstructured":"Abed A, Akrout B, Amous I (2024) Deep learning-based few-shot person re-identification from top-view rgb and depth images. Neural Comput Appl 36(31):19365\u201319382","journal-title":"Neural Comput Appl"},{"issue":"3","key":"7139_CR47","doi-asserted-by":"publisher","first-page":"257","DOI":"10.1109\/JPROC.2023.3238524","volume":"111","author":"Z Zou","year":"2023","unstructured":"Zou Z, Chen K, Shi Z, Guo Y, Ye J (2023) Object detection in 20 years: A survey. Proc IEEE 111(3):257\u2013276","journal-title":"Proc IEEE"},{"key":"7139_CR48","doi-asserted-by":"crossref","unstructured":"Zhang C, Lam KM, Liu T, Chan YL, Wang Q (2024) Structured adversarial self-supervised learning for robust object detection in remote sensing images. IEEE Trans Geosci Remote Sens","DOI":"10.1109\/TGRS.2024.3375398"},{"key":"7139_CR49","unstructured":"Yao Z, Ai J, Li B, Zhang C (2021) Efficient detr: improving end-to-end object detector with dense prior. arXiv:2104.01318"},{"key":"7139_CR50","unstructured":"Chen Q, Su X, Zhang X, Wang J, Chen J, Shen Y, Han C, Chen Z, Xu W, Li F et al (2024) Lw-detr: A transformer replacement to yolo for real-time detection. arXiv:2406.03459"},{"key":"7139_CR51","unstructured":"Zhao T, Liu P, He X, Zhang L, Lee K (2024) Real-time transformer-based open-vocabulary detection with efficient fusion head. arXiv preprint arXiv:2403.06892"},{"key":"7139_CR52","unstructured":"Grewal MS, Andrews AP (2014) Kalman filtering: Theory and Practice with MATLAB. John Wiley & Sons"},{"issue":"6","key":"7139_CR53","doi-asserted-by":"publisher","first-page":"1592","DOI":"10.5465\/amj.2011.0932","volume":"57","author":"WK Smith","year":"2014","unstructured":"Smith WK (2014) Dynamic decision making: A model of senior leaders managing strategic paradoxes. Acad Manag J 57(6):1592\u20131623","journal-title":"Acad Manag J"},{"key":"7139_CR54","doi-asserted-by":"crossref","unstructured":"Wang K, Liew JH, Zou Y, Zhou D, Feng J (2019) Panet: Few-shot image semantic segmentation with prototype alignment. In: proceedings of the IEEE\/CVF international conference on computer vision, pp 9197\u20139206","DOI":"10.1109\/ICCV.2019.00929"},{"key":"7139_CR55","doi-asserted-by":"crossref","unstructured":"Chang MF, Lambert J, Sangkloy P, Singh J, Bak S, Hartnett A, Wang D, Carr P, Lucey S, Ramanan D et al (2019) Argoverse: 3d tracking and forecasting with rich maps. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 8748\u20138757","DOI":"10.1109\/CVPR.2019.00895"},{"key":"7139_CR56","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 770\u2013778","DOI":"10.1109\/CVPR.2016.90"},{"key":"7139_CR57","doi-asserted-by":"crossref","unstructured":"Meng D, Chen X, Fan Z, Zeng G, Li H, Yuan Y, Sun L, Wang J (2021) Conditional detr for fast training convergence. In: 2021 IEEE\/CVF International Conference on Computer Vision (ICCV), IEEE, pp 3631\u20133640","DOI":"10.1109\/ICCV48922.2021.00363"},{"key":"7139_CR58","doi-asserted-by":"crossref","unstructured":"Ye M, Ke L, Li S, Tai YW, Tang CK, Danelljan M, Yu F (2023) Cascade-detr: Delving into high-quality universal object detection. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 6704\u20136714","DOI":"10.1109\/ICCV51070.2023.00617"},{"key":"7139_CR59","unstructured":"Lin J, Mao X, Chen Y, Xu L, He Y, Xue H (2022) D$$\\hat{~}$$2etr: Decoder-only detr with computationally efficient cross-scale attention. arXiv:2203.00860"},{"key":"7139_CR60","doi-asserted-by":"crossref","unstructured":"Gao P, Zheng M, Wang X, Dai J, Li H (2021) Fast convergence of detr with spatially modulated co-attention. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 3621\u20133630","DOI":"10.1109\/ICCV48922.2021.00360"},{"issue":"16","key":"7139_CR61","doi-asserted-by":"publisher","first-page":"7107","DOI":"10.1109\/JSEN.2019.2913281","volume":"19","author":"X Liang","year":"2019","unstructured":"Liang X, Hu P, Zhang L, Sun J, Yin G (2019) Mcfnet: Multi-layer concatenation fusion network for medical images fusion. IEEE Sens J 19(16):7107\u20137119","journal-title":"IEEE Sens J"},{"key":"7139_CR62","doi-asserted-by":"crossref","unstructured":"Ren M, Zemel RS (2016) End-to-end instance segmentation with recurrent attention. 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR) pp 293\u2013301. https:\/\/api.semanticscholar.org\/CorpusID:206595795","DOI":"10.1109\/CVPR.2017.39"},{"key":"7139_CR63","doi-asserted-by":"publisher","first-page":"103911","DOI":"10.1016\/j.imavis.2020.103911","volume":"97","author":"S Wu","year":"2020","unstructured":"Wu S, Li X, Wang X (2020) Iou-aware single-stage object detector for accurate localization. Image Vis Comput 97:103911","journal-title":"Image Vis Comput"},{"key":"7139_CR64","doi-asserted-by":"publisher","first-page":"107816","DOI":"10.1016\/j.patcog.2021.107816","volume":"112","author":"L Zhu","year":"2021","unstructured":"Zhu L, Xie Z, Liu L, Tao B, Tao W (2021) Iou-uniform r-cnn: Breaking through the limitations of rpn. Pattern Recogn 112:107816","journal-title":"Pattern Recogn"},{"key":"7139_CR65","doi-asserted-by":"crossref","unstructured":"Tan M, Pang R, Le QV (2020) Efficientdet: Scalable and efficient object detection. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 10781\u201310790","DOI":"10.1109\/CVPR42600.2020.01079"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-026-07139-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-026-07139-8","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-026-07139-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,3,31]],"date-time":"2026-03-31T05:01:08Z","timestamp":1774933268000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-026-07139-8"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,2]]},"references-count":65,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2026,2]]}},"alternative-id":["7139"],"URL":"https:\/\/doi.org\/10.1007\/s10489-026-07139-8","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,2]]},"assertion":[{"value":"18 November 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 February 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"14 February 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"We would like to note that in the manuscript entitled \u201cStream-DINO: Exploring DETR-Based Online Object Detection with Streaming Perception\u201d, no conflict of interest exists in the submission of this manuscript, and the manuscript is approved by all authors for publication.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}},{"value":"We confirm that the used data does not include any Ethical consent.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical and informed consent"}}],"article-number":"90"}}