{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,23]],"date-time":"2026-07-23T09:24:45Z","timestamp":1784798685883,"version":"3.55.0"},"reference-count":50,"publisher":"Springer Science and Business Media LLC","issue":"8","license":[{"start":{"date-parts":[[2025,5,31]],"date-time":"2025-05-31T00:00:00Z","timestamp":1748649600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,5,31]],"date-time":"2025-05-31T00:00:00Z","timestamp":1748649600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["61962007, 62266009"],"award-info":[{"award-number":["61962007, 62266009"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["61962007, 62266009"],"award-info":[{"award-number":["61962007, 62266009"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Guangxi Key Laboratory of Big Data in Finance and Economics","award":["FEDOP2022A06"],"award-info":[{"award-number":["FEDOP2022A06"]}]},{"name":"Guangxi Key Laboratory of Big Data in Finance and Economics","award":["FEDOP2022A06"],"award-info":[{"award-number":["FEDOP2022A06"]}]},{"name":"Fundamental Research Capacity Improvement Program for Young and Middle-Aged Teachers in Guangxi Higher Education Institutions","award":["2022KY1477"],"award-info":[{"award-number":["2022KY1477"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["J Supercomput"],"DOI":"10.1007\/s11227-025-07381-w","type":"journal-article","created":{"date-parts":[[2025,5,31]],"date-time":"2025-05-31T06:12:28Z","timestamp":1748671948000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":4,"title":["Multi-object tracking in surveillance scenarios based on joint detection and multi-feature fusion"],"prefix":"10.1007","volume":"81","author":[{"given":"Runxing","family":"Zhao","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jianxin","family":"Gong","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zhiwen","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,5,31]]},"reference":[{"key":"7381_CR1","doi-asserted-by":"publisher","first-page":"6400","DOI":"10.1007\/s10489-021-02293-7","volume":"51","author":"SK Pal","year":"2021","unstructured":"Pal SK, Pramanik A, Maiti J, Mitra P (2021) Deep learning in multi-object detection and tracking: state of the art. Appl Intell 51:6400\u20136429. https:\/\/doi.org\/10.1007\/s10489-021-02293-7","journal-title":"Appl Intell"},{"key":"7381_CR2","doi-asserted-by":"publisher","DOI":"10.1016\/j.artint.2020.103448","volume":"293","author":"W Luo","year":"2021","unstructured":"Luo W, Xing J, Milan A et al (2021) Multiple object tracking: a literature review. Artif Intell 293:103448. https:\/\/doi.org\/10.1016\/j.artint.2020.103448","journal-title":"Artif Intell"},{"key":"7381_CR3","doi-asserted-by":"publisher","first-page":"107","DOI":"10.1007\/978-3-030-58621-8_7","volume-title":"Computer Vision \u2013 ECCV 2020","author":"Z Wang","year":"2020","unstructured":"Wang Z, Zheng L, Liu Y et al (2020) Towards real-time multi-object tracking. In: Vedaldi A, Bischof H, Brox T, Frahm J-M (eds) Computer Vision \u2013 ECCV 2020. Springer International Publishing, Cham, pp 107\u2013122"},{"key":"7381_CR4","doi-asserted-by":"publisher","first-page":"3069","DOI":"10.1007\/s11263-021-01513-4","volume":"129","author":"Y Zhang","year":"2021","unstructured":"Zhang Y, Wang C, Wang X et al (2021) FairMOT: on the fairness of detection and re-identification in multiple object tracking. Int J Comput Vis 129:3069\u20133087. https:\/\/doi.org\/10.1007\/s11263-021-01513-4","journal-title":"Int J Comput Vis"},{"key":"7381_CR5","doi-asserted-by":"crossref","unstructured":"Duan K, Bai S, Xie L, et al (2019) CenterNet: keypoint triplets for object detection. In: 2019 IEEE\/CVF International Conference on Computer Vision (ICCV). pp 6568\u20136577","DOI":"10.1109\/ICCV.2019.00667"},{"key":"7381_CR6","unstructured":"Ge Z, Liu S, Wang F, et al (2021) YOLOX: Exceeding YOLO series in 2021"},{"key":"7381_CR7","doi-asserted-by":"publisher","first-page":"474","DOI":"10.1007\/978-3-030-58548-8_28","volume-title":"Computer Vision \u2013 ECCV 2020","author":"X Zhou","year":"2020","unstructured":"Zhou X, Koltun V, Kr\u00e4henb\u00fchl P (2020) Tracking objects as points. In: Vedaldi A, Bischof H, Brox T, Frahm J-M (eds) Computer Vision \u2013 ECCV 2020. Springer International Publishing, Cham, pp 474\u2013490"},{"key":"7381_CR8","doi-asserted-by":"publisher","first-page":"3182","DOI":"10.1109\/TIP.2022.3165376","volume":"31","author":"C Liang","year":"2022","unstructured":"Liang C, Zhang Z, Zhou X et al (2022) Rethinking the competition between detection and reid in multiobject tracking. IEEE Trans Image Process 31:3182\u20133196. https:\/\/doi.org\/10.1109\/TIP.2022.3165376","journal-title":"IEEE Trans Image Process"},{"key":"7381_CR9","doi-asserted-by":"crossref","unstructured":"Bewley A, Ge Z, Ott L, et al (2016) Simple online and realtime tracking. In: 2016 IEEE International Conference on Image Processing (ICIP). pp 3464\u20133468","DOI":"10.1109\/ICIP.2016.7533003"},{"key":"7381_CR10","doi-asserted-by":"crossref","unstructured":"Bochinski E, Eiselein V, Sikora T (2017) High-speed tracking-by-detection without using image information. In: 2017 14th IEEE International Conference on Advanced Video and Signal Based Surveillance (AVSS). pp 1\u20136","DOI":"10.1109\/AVSS.2017.8078516"},{"key":"7381_CR11","doi-asserted-by":"crossref","unstructured":"Wojke N, Bewley A, Paulus D (2017) Simple Online and Realtime Tracking with a Deep Association Metric","DOI":"10.1109\/ICIP.2017.8296962"},{"key":"7381_CR12","doi-asserted-by":"crossref","unstructured":"Tang S, Andriluka M, Andres B, Schiele B (2017) Multiple people tracking by lifted multicut and person re-identification. In: 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR). pp 3701\u20133710","DOI":"10.1109\/CVPR.2017.394"},{"key":"7381_CR13","doi-asserted-by":"publisher","first-page":"7077","DOI":"10.1007\/s11042-018-6467-6","volume":"78","author":"N Mahmoudi","year":"2019","unstructured":"Mahmoudi N, Ahadi SM, Rahmati M (2019) Multi-target tracking using CNN-based features: CNNMTT. Multimed Tools Appl 78:7077\u20137096. https:\/\/doi.org\/10.1007\/s11042-018-6467-6","journal-title":"Multimed Tools Appl"},{"key":"7381_CR14","doi-asserted-by":"crossref","unstructured":"Fang K, Xiang Y, Li X, Savarese S (2018) Recurrent autoregressive networks for online multi-object tracking","DOI":"10.1109\/WACV.2018.00057"},{"key":"7381_CR15","doi-asserted-by":"crossref","unstructured":"Luo H, Gu Y, Liao X, et al (2019) Bag of tricks and a strong baseline for deep person re-identification","DOI":"10.1109\/CVPRW.2019.00190"},{"key":"7381_CR16","doi-asserted-by":"crossref","unstructured":"Zhou Z, Xing J, Zhang M, Hu W (2018) Online multi-target tracking with tensor-based high-order graph matching. In: 2018 24th International Conference on Pattern Recognition (ICPR). pp 1809\u20131814","DOI":"10.1109\/ICPR.2018.8545450"},{"key":"7381_CR17","doi-asserted-by":"publisher","DOI":"10.1007\/s11042-023-15862-4","author":"L Huang","year":"2023","unstructured":"Huang L, Wang Z, Fu X (2023) Pedestrian detection using RetinaNet with multi-branch structure and double pooling attention mechanism. Multimed Tools Appl. https:\/\/doi.org\/10.1007\/s11042-023-15862-4","journal-title":"Multimed Tools Appl"},{"key":"7381_CR18","doi-asserted-by":"crossref","unstructured":"Zhang Y, Sun P, Jiang Y, et al (2022) ByteTrack: multi-object tracking by associating every detection box","DOI":"10.1007\/978-3-031-20047-2_1"},{"key":"7381_CR19","unstructured":"Aharon N, Orfaig R, Bobrovsky B-Z (2022) BoT-SORT: robust associations multi-pedestrian tracking"},{"key":"7381_CR20","doi-asserted-by":"publisher","first-page":"1462","DOI":"10.1109\/TMM.2023.3234822","volume":"25","author":"Z Liu","year":"2023","unstructured":"Liu Z, Shang Y, Li T et al (2023) Robust multi-drone multi-target tracking to resolve target occlusion: a benchmark. IEEE Trans Multimedia 25:1462\u20131476. https:\/\/doi.org\/10.1109\/TMM.2023.3234822","journal-title":"IEEE Trans Multimedia"},{"key":"7381_CR21","doi-asserted-by":"publisher","first-page":"1256","DOI":"10.1109\/TMM.2022.3140919","volume":"25","author":"G Wang","year":"2023","unstructured":"Wang G, Wang Y, Gu R et al (2023) Split and connect: a universal tracklet booster for multi-object tracking. IEEE Trans Multimed 25:1256\u20131268. https:\/\/doi.org\/10.1109\/TMM.2022.3140919","journal-title":"IEEE Trans Multimed"},{"key":"7381_CR22","doi-asserted-by":"publisher","first-page":"39655","DOI":"10.1007\/s11042-022-13058-w","volume":"81","author":"Z Wang","year":"2022","unstructured":"Wang Z, Feng J, Zhang Y (2022) Pedestrian detection in infrared image based on depth transfer learning. Multimed Tools Appl 81:39655\u201339674. https:\/\/doi.org\/10.1007\/s11042-022-13058-w","journal-title":"Multimed Tools Appl"},{"key":"7381_CR23","doi-asserted-by":"crossref","unstructured":"Feichtenhofer C, Pinz A, Zisserman A (2018) Detect to track and track to detect","DOI":"10.1109\/ICCV.2017.330"},{"key":"7381_CR24","doi-asserted-by":"crossref","unstructured":"Chen L, Ai H, Zhuang Z, Shang C (2018) Real-time multiple people tracking with deeply learned candidate selection and person re-identification. In: 2018 IEEE International Conference on Multimedia and Expo (ICME). pp 1\u20136","DOI":"10.1109\/ICME.2018.8486597"},{"key":"7381_CR25","doi-asserted-by":"crossref","unstructured":"Bergmann P, Meinhardt T, Leal-Taixe L (2019) Tracking without bells and whistles. In: 2019 IEEE\/CVF International Conference on Computer Vision (ICCV). pp 941\u2013951","DOI":"10.1109\/ICCV.2019.00103"},{"key":"7381_CR26","unstructured":"Zhang J, Zhou S, Chang X, et al (2020) Multiple object tracking by flowing and fusing"},{"key":"7381_CR27","doi-asserted-by":"crossref","unstructured":"Tokmakov P, Li J, Burgard W, Gaidon A (2021) Learning to track with object permanence","DOI":"10.1109\/ICCV48922.2021.01068"},{"key":"7381_CR28","doi-asserted-by":"crossref","unstructured":"Lu Z, Rathod V, Votel R, Huang J (2020) RetinaTrack: online single stage joint detection and tracking","DOI":"10.1109\/CVPR42600.2020.01468"},{"key":"7381_CR29","doi-asserted-by":"crossref","unstructured":"Wang Q, Zheng Y, Pan P, Xu Y (2021) Multiple object tracking with correlation learning","DOI":"10.1109\/CVPR46437.2021.00387"},{"key":"7381_CR30","doi-asserted-by":"publisher","DOI":"10.1016\/j.eswa.2023.120271","volume":"226","author":"R Zhao","year":"2023","unstructured":"Zhao R, Wang Z, Guo W, Zhang C (2023) Multi-scene image enhancement based on multi-channel illumination estimation. Expert Syst Appl 226:120271. https:\/\/doi.org\/10.1016\/j.eswa.2023.120271","journal-title":"Expert Syst Appl"},{"key":"7381_CR31","doi-asserted-by":"crossref","unstructured":"Wu J, Cao J, Song L, et al (2021) Track to detect and segment: an online multi-object tracker. In: 2021 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). pp 12347\u201312356","DOI":"10.1109\/CVPR46437.2021.01217"},{"key":"7381_CR32","doi-asserted-by":"publisher","first-page":"4445","DOI":"10.1109\/TMM.2023.3323852","volume":"26","author":"J Zhang","year":"2024","unstructured":"Zhang J, Wang M, Jiang H et al (2024) STAT: multi-object tracking based on spatio-temporal topological constraints. IEEE Trans Multimed 26:4445\u20134457. https:\/\/doi.org\/10.1109\/TMM.2023.3323852","journal-title":"IEEE Trans Multimed"},{"key":"7381_CR33","doi-asserted-by":"publisher","DOI":"10.1016\/j.eswa.2024.123581","volume":"249","author":"S Ma","year":"2024","unstructured":"Ma S, Duan S, Hou Z et al (2024) Multi-object tracking algorithm based on interactive attention network and adaptive trajectory reconnection. Expert Syst Appl 249:123581. https:\/\/doi.org\/10.1016\/j.eswa.2024.123581","journal-title":"Expert Syst Appl"},{"key":"7381_CR34","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2022.105770","volume":"119","author":"C-Y Tsai","year":"2023","unstructured":"Tsai C-Y, Shen G-Y, Nisar H (2023) Swin-JDE: joint detection and embedding multi-object tracking in crowded scenes based on swin-transformer. Eng Appl Artif Intell 119:105770. https:\/\/doi.org\/10.1016\/j.engappai.2022.105770","journal-title":"Eng Appl Artif Intell"},{"key":"7381_CR35","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.127328","volume":"575","author":"S Chan","year":"2024","unstructured":"Chan S, Qiu C, Wu D et al (2024) Fusion detection and ReID embedding with hybrid attention for multi-object tracking. Neurocomputing 575:127328. https:\/\/doi.org\/10.1016\/j.neucom.2024.127328","journal-title":"Neurocomputing"},{"key":"7381_CR36","doi-asserted-by":"publisher","DOI":"10.1007\/s40747-024-01426-y","author":"X Feng","year":"2024","unstructured":"Feng X, Jiao X, Wang S et al (2024) SCGTracker: object feature embedding enhancement based on graph attention networks for multi-object tracking. Complex Intell Syst. https:\/\/doi.org\/10.1007\/s40747-024-01426-y","journal-title":"Complex Intell Syst"},{"key":"7381_CR37","doi-asserted-by":"crossref","unstructured":"Wang C-Y, Yeh I-H, Liao H-YM (2024) YOLOv9: Learning what you want to learn using programmable gradient information","DOI":"10.1007\/978-3-031-72751-1_1"},{"key":"7381_CR38","doi-asserted-by":"publisher","first-page":"1489","DOI":"10.1109\/TPAMI.2022.3164083","volume":"45","author":"Y Li","year":"2023","unstructured":"Li Y, Yao T, Pan Y, Mei T (2023) Contextual transformer networks for visual recognition. IEEE Trans Pattern Anal Mach Intell 45:1489\u20131500. https:\/\/doi.org\/10.1109\/TPAMI.2022.3164083","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"7381_CR39","doi-asserted-by":"crossref","unstructured":"Li Y, Hu J, Wen Y, et al (2023) Rethinking vision transformers for MobileNet size and speed","DOI":"10.1109\/ICCV51070.2023.01549"},{"key":"7381_CR40","doi-asserted-by":"crossref","unstructured":"Zhou K, Yang Y, Cavallaro A, Xiang T (2019) Omni-scale feature learning for person re-identification","DOI":"10.1109\/ICCV.2019.00380"},{"key":"7381_CR41","unstructured":"Gevorgyan Z (2022) SIoU Loss: more powerful learning for bounding box regression"},{"key":"7381_CR42","doi-asserted-by":"crossref","unstructured":"Cipolla R, Gal Y, Kendall A (2018) Multi-task learning using uncertainty to weigh losses for scene geometry and semantics. In: 2018 IEEE\/CVF Conference on Computer Vision and Pattern Recognition. pp 7482\u20137491","DOI":"10.1109\/CVPR.2018.00781"},{"key":"7381_CR43","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1155\/2008\/246309","volume":"2008","author":"K Bernardin","year":"2008","unstructured":"Bernardin K, Stiefelhagen R (2008) Evaluating multiple object tracking performance: the clear mot metrics. J Image Video Proc 2008:1\u201310. https:\/\/doi.org\/10.1155\/2008\/246309","journal-title":"J Image Video Proc"},{"key":"7381_CR44","doi-asserted-by":"publisher","first-page":"17","DOI":"10.1007\/978-3-319-48881-3_2","volume-title":"Computer Vision \u2013 ECCV 2016 Workshops","author":"E Ristani","year":"2016","unstructured":"Ristani E, Solera F, Zou R et al (2016) Performance measures and a data set for multi-target, multi-camera tracking. In: Hua G, J\u00e9gou H (eds) Computer Vision \u2013 ECCV 2016 Workshops. Springer International Publishing, Cham, pp 17\u201335"},{"key":"7381_CR45","doi-asserted-by":"publisher","first-page":"548","DOI":"10.1007\/s11263-020-01375-2","volume":"129","author":"JA LuitenOs\u0306ep","year":"2021","unstructured":"LuitenOs\u0306ep JA, Dendorfer P et al (2021) HOTA: a higher order metric for evaluating multi-object tracking. Int J Comput Vis 129:548\u2013578. https:\/\/doi.org\/10.1007\/s11263-020-01375-2","journal-title":"Int J Comput Vis"},{"key":"7381_CR46","doi-asserted-by":"crossref","unstructured":"Wang C-Y, Bochkovskiy A, Liao H-YM (2022) YOLOv7: Trainable bag-of-freebies sets new state-of-the-art for real-time object detectors","DOI":"10.1109\/CVPR52729.2023.00721"},{"key":"7381_CR47","doi-asserted-by":"crossref","unstructured":"Wang Y, Kitani K, Weng X (2021) Joint object detection and multi-object tracking with graph neural networks. In: 2021 IEEE International Conference on Robotics and Automation (ICRA). pp 13708\u201313715","DOI":"10.1109\/ICRA48506.2021.9561110"},{"key":"7381_CR48","doi-asserted-by":"publisher","first-page":"333","DOI":"10.1016\/j.neucom.2022.01.008","volume":"483","author":"Q Liu","year":"2022","unstructured":"Liu Q, Chen D, Chu Q et al (2022) Online multi-object tracking with unsupervised re-identification learning and occlusion estimation. Neurocomputing 483:333\u2013347. https:\/\/doi.org\/10.1016\/j.neucom.2022.01.008","journal-title":"Neurocomputing"},{"key":"7381_CR49","doi-asserted-by":"publisher","first-page":"4378","DOI":"10.1109\/TIP.2023.3298538","volume":"32","author":"S-H Lee","year":"2023","unstructured":"Lee S-H, Park D-H, Bae S-H (2023) Decode-MOT: how can we hurdle frames to go beyond tracking-by-detection? IEEE Trans Image Process 32:4378\u20134392. https:\/\/doi.org\/10.1109\/TIP.2023.3298538","journal-title":"IEEE Trans Image Process"},{"key":"7381_CR50","doi-asserted-by":"publisher","first-page":"5115","DOI":"10.1007\/s40747-023-01009-3","volume":"9","author":"J Cao","year":"2023","unstructured":"Cao J, Zhang J, Li B et al (2023) RetinaMOT: rethinking anchor-free YOLOv5 for online multiple object tracking. Complex Intell Syst 9:5115\u20135133. https:\/\/doi.org\/10.1007\/s40747-023-01009-3","journal-title":"Complex Intell Syst"}],"container-title":["The Journal of Supercomputing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11227-025-07381-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11227-025-07381-w\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11227-025-07381-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,5,31]],"date-time":"2025-05-31T06:12:36Z","timestamp":1748671956000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11227-025-07381-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,5,31]]},"references-count":50,"journal-issue":{"issue":"8","published-online":{"date-parts":[[2025,6]]}},"alternative-id":["7381"],"URL":"https:\/\/doi.org\/10.1007\/s11227-025-07381-w","relation":{},"ISSN":["1573-0484"],"issn-type":[{"value":"1573-0484","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,5,31]]},"assertion":[{"value":"2 May 2025","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"31 May 2025","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no competing interests.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"937"}}