{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T17:03:16Z","timestamp":1784566996355,"version":"3.55.0"},"reference-count":91,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2026,5,15]],"date-time":"2026-05-15T00:00:00Z","timestamp":1778803200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,15]],"date-time":"2026-05-15T00:00:00Z","timestamp":1778803200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s11263-026-02858-4","type":"journal-article","created":{"date-parts":[[2026,5,15]],"date-time":"2026-05-15T16:11:26Z","timestamp":1778861486000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["NOOUGAT \n              \n             Towards Unified Online and Offline Multi-Object Tracking"],"prefix":"10.1007","volume":"134","author":[{"ORCID":"https:\/\/orcid.org\/0009-0008-9522-8214","authenticated-orcid":false,"given":"Benjamin","family":"Missaoui","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Orcun","family":"Cetintas","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Guillem","family":"Bras\u00f3","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Tim","family":"Meinhardt","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Laura","family":"Leal-Taix\u00e9","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,15]]},"reference":[{"key":"2858_CR1","unstructured":"Aharon, N., Orfaig, R., & Bobrovsky, B. Z. (2022). BoT-SORT: Robust Associations Multi-Pedestrian Tracking. arXiv:2206.14651."},{"issue":"9","key":"2858_CR2","doi-asserted-by":"publisher","first-page":"1806","DOI":"10.1109\/TPAMI.2011.21","volume":"33","author":"J Berclaz","year":"2011","unstructured":"Berclaz, J., Fleuret, F., Turetken, E., & Fua, P. (2011). Multiple object tracking using k-shortest paths optimization. IEEE TPAMI., 33(9), 1806\u20131819.","journal-title":"IEEE TPAMI."},{"key":"2858_CR3","doi-asserted-by":"crossref","unstructured":"Bergmann, P., Meinhardt, T., & Leal-Taix\u00e9, L. (2019). Tracking without bells and whistles. In: The IEEE International Conference on Computer Vision (ICCV).","DOI":"10.1109\/ICCV.2019.00103"},{"key":"2858_CR4","doi-asserted-by":"crossref","unstructured":"Bewley, A., Ge, Z., Ott, L., Ramos, F., & Upcroft, B. (2016). Simple online and realtime tracking. In: 2016 IEEE International Conference on Image Processing (ICIP);. p. 3464\u20133468.","DOI":"10.1109\/ICIP.2016.7533003"},{"key":"2858_CR5","doi-asserted-by":"publisher","DOI":"10.1007\/s11263-022-01678-6","author":"G Bras\u00f3","year":"2022","unstructured":"Bras\u00f3, G., Cetintas, O., & Leal-Taix\u00e9, L. (2022). Multi-Object Tracking and Segmentation Via Neural Message Passing. International Journal of Computer Vision. https:\/\/doi.org\/10.1007\/s11263-022-01678-6","journal-title":"International Journal of Computer Vision."},{"issue":"12","key":"2858_CR6","doi-asserted-by":"publisher","first-page":"3035","DOI":"10.1007\/s11263-022-01678-6","volume":"130","author":"G Bras\u00f3","year":"2022","unstructured":"Bras\u00f3, G., Cetintas, O., & Leal-Taix\u00e9, L. (2022). Multi-Object Tracking and Segmentation Via Neural Message Passing. International Journal of Computer Vision., 130(12), 3035\u20133053.","journal-title":"International Journal of Computer Vision."},{"key":"2858_CR7","doi-asserted-by":"crossref","unstructured":"Bras\u00f3, G., & Leal-Taix\u00e9, L. (2020). Learning a Neural Solver for Multiple Object Tracking. In: The IEEE Conference on Computer Vision and Pattern Recognition (CVPR).","DOI":"10.1109\/CVPR42600.2020.00628"},{"key":"2858_CR8","doi-asserted-by":"crossref","unstructured":"Butt, A., & Collins, R. (2013). Multi-target Tracking by Lagrangian Relaxation to Min-Cost Network Flow. CVPR.","DOI":"10.1109\/CVPR.2013.241"},{"key":"2858_CR9","doi-asserted-by":"crossref","unstructured":"Cai, J., Xu, M., Li, W., Xiong, Y., Xia, W., Tu, Z., & Soatto, S. (2022). MeMOT: Multi-Object Tracking with Memory. In: 2022 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). p. 8080\u20138090.","DOI":"10.1109\/CVPR52688.2022.00792"},{"key":"2858_CR10","doi-asserted-by":"crossref","unstructured":"Cao, J., Pang, J., Weng, X., Khirodkar, R., & Kitani, K. (2023). Observation-centric sort: Rethinking sort for robust multi-object tracking. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. p. 9686\u20139696.","DOI":"10.1109\/CVPR52729.2023.00934"},{"key":"2858_CR11","doi-asserted-by":"publisher","first-page":"743","DOI":"10.1109\/TIP.2025.3526066","volume":"34","author":"X Cao","year":"2025","unstructured":"Cao, X., Zheng, Y., Yao, Y., Qin, H., Cao, X., & Guo, S. (2025). TOPIC: A Parallel Association Paradigm for Multi-Object Tracking Under Complex Motions and Diverse Scenes. IEEE Transactions on Image Processing., 34, 743\u2013758. https:\/\/doi.org\/10.1109\/TIP.2025.3526066","journal-title":"IEEE Transactions on Image Processing."},{"key":"2858_CR12","doi-asserted-by":"publisher","unstructured":"Carion, N., Massa, F., Synnaeve, G., Usunier, N., Kirillov, A., & Zagoruyko, S. (2020). End-to-End Object Detection with Transformers. In: Computer Vision \u2013 ECCV 2020: 16th European Conference, Glasgow, UK, August 23\u201328, 2020, Proceedings, Part I Berlin, Heidelberg: Springer-Verlag. p. 213\u201322https:\/\/doi.org\/10.1007\/978-3-030-58452-8_13.","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"2858_CR13","doi-asserted-by":"crossref","unstructured":"Cetintas, O., Bras\u00f3, G., & Leal-Taix\u00e9, L. (2023). Unifying Short and Long-Term Tracking With Graph Hierarchies. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). p. 22877\u201322887.","DOI":"10.1109\/CVPR52729.2023.02191"},{"key":"2858_CR14","doi-asserted-by":"crossref","unstructured":"Cetintas, O., Meinhardt, T., Bras\u00f3, G., & Leal-Taix\u00e9, L. (2024). SPAMming Labels: Efficient Annotations for the Trackers of Tomorrow. In: European Conference on Computer Vision (ECCV).","DOI":"10.1007\/978-3-031-73254-6_22"},{"key":"2858_CR15","doi-asserted-by":"crossref","unstructured":"Cui, Y., Zeng, C., Zhao, X., Yang, Y., Wu, G., & Wang, L. (2023). SportsMOT: A Large Multi-Object Tracking Dataset in Multiple Sports Scenes. arXiv preprint arXiv:2304.05170.","DOI":"10.1109\/ICCV51070.2023.00910"},{"key":"2858_CR16","doi-asserted-by":"crossref","unstructured":"Dai P, Weng, R., Choi, W., Zhang, C., He, Z., & Ding, W. (2021). Learning a Proposal Classifier for Multiple Object Tracking. In: CVPR. p. 2443\u20132452.","DOI":"10.1109\/CVPR46437.2021.00247"},{"key":"2858_CR17","unstructured":"Dendorfer, P., Rezatofighi, H., Milan, A., Shi, J., Cremers, D., Reid, I., Roth, S., Schindler, K., & Leal-Taix\u00e9, L. (2020). MOT20: A benchmark for multi object tracking in crowded scenes. arXiv: 2003.09003."},{"key":"2858_CR18","doi-asserted-by":"crossref","unstructured":"Dendorfer, P., O\u0161ep, A., Milan, A., Schindler, K., Cremers, D., Reid, I., Roth, S., & Leal-Taix\u00e9, L. (2020). MOTChallenge: A Benchmark for Single-Camera Multiple Target Tracking. arXiv:2010.07548.","DOI":"10.1007\/s11263-020-01393-0"},{"key":"2858_CR19","doi-asserted-by":"crossref","unstructured":"Dendorfer, P., Yugay, V., Osep, A., & Leal-Taix\u00e9, L. (2022). Quo Vadis: Is Trajectory Forecasting the Key Towards Long-Term Multi-Object Tracking? In: Oh AH, Agarwal A, Belgrave D, Cho K, editors. Advances in Neural Information Processing Systems. https:\/\/openreview.net\/forum?id=3r0yLLCo4fF.","DOI":"10.52202\/068431-1139"},{"key":"2858_CR20","doi-asserted-by":"crossref","unstructured":"Ding, S., Schneider, L., Cordts, M., & Gall, J. (2024). ADA-Track++: End-to-End Multi-Camera 3D Multi-Object Tracking with Alternating Detection and Association.","DOI":"10.1109\/CVPR52733.2024.01438"},{"key":"2858_CR21","doi-asserted-by":"crossref","unstructured":"Du, Y., Zhao, Z., Song, Y., Zhao, Y., Su F., Gong, T., & Meng, H. (2023). Strongsort: Make deepsort great again. IEEE Transactions on Multimedia.","DOI":"10.1109\/TMM.2023.3240881"},{"key":"2858_CR22","doi-asserted-by":"crossref","unstructured":"Fu, D., Chen, D., Bao, J., Yang, H., Yuan, L., Zhang, L., Li, H., & Chen, D. (2021). Unsupervised Pre-training for Person Re-identification. Proceedings of the IEEE conference on computer vision and pattern recognition.","DOI":"10.1109\/CVPR46437.2021.01451"},{"key":"2858_CR23","doi-asserted-by":"crossref","unstructured":"Gao, R., Qi, J., & Wang, L. (2025). Multiple Object Tracking as ID Prediction. arXiv:2403.16848.","DOI":"10.1109\/CVPR52734.2025.02596"},{"key":"2858_CR24","doi-asserted-by":"crossref","unstructured":"Gao, R., & Wang, L. (2023). MeMOTR: Long-Term Memory-Augmented Transformer for Multi-Object Tracking. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV). p. 9901\u20139910.","DOI":"10.1109\/ICCV51070.2023.00908"},{"key":"2858_CR25","doi-asserted-by":"publisher","unstructured":"Gao, Y., Xu, H., Li, J., Wang, N., & Gao, X. (2024). Multi-scene generalized trajectory global graph solver with composite nodes for multiple object tracking. In: Proceedings of the Thirty-Eighth AAAI Conference on Artificial Intelligence and Thirty-Sixth Conference on Innovative Applications of Artificial Intelligence and Fourteenth Symposium on Educational Advances in Artificial Intelligence AAAI\u201924\/IAAI\u201924\/EAAI\u201924, AAAI Press.https:\/\/doi.org\/10.1609\/aaai.v38i3.27953.","DOI":"10.1609\/aaai.v38i3.27953"},{"key":"2858_CR26","unstructured":"Ge, Z., Liu, S., Wang, F., Li, Z., & Sun, J. (2021). YOLOX: Exceeding YOLO Series in 2021. arXiv preprint arXiv:2107.08430."},{"key":"2858_CR27","doi-asserted-by":"crossref","unstructured":"Han, X., Oishi, N., Tian, Y., Ucurum, E., Young, R., Chatwin, C., & Birch, P. (2024). ETTrack: Enhanced Temporal Motion Predictor for Multi-Object Tracking. arXiv:2405.15755.","DOI":"10.1007\/s10489-024-05866-4"},{"key":"2858_CR28","doi-asserted-by":"crossref","unstructured":"He, J., Huang, Z., Wang, N., & Zhang, Z. (2021). Learnable graph matching: Incorporating graph partitioning with deep feature learning for multiple object tracking. In: CVPR. p. 5299\u20135309.","DOI":"10.1109\/CVPR46437.2021.00526"},{"key":"2858_CR29","doi-asserted-by":"crossref","unstructured":"He, J., Huang, Z., Wang, N., & Zhang, Z. (2021). Learnable Graph Matching: Incorporating Graph Partitioning With Deep Feature Learning for Multiple Object Tracking. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). p. 5299\u20135309.","DOI":"10.1109\/CVPR46437.2021.00526"},{"key":"2858_CR30","unstructured":"He, L., Liao, X., Liu, W., Liu, X., Cheng, P., & Mei, T. (2020). FastReID: A Pytorch Toolbox for General Instance Re-identification. arXiv preprint arXiv:2006.02631."},{"key":"2858_CR31","doi-asserted-by":"crossref","unstructured":"He, S., Luo, H., Wang, P., Wang, F., Li, H., & Jiang, W. (2021). TransReID: Transformer-Based Object Re-Identification. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV); . p. 15013\u201315022.","DOI":"10.1109\/ICCV48922.2021.01474"},{"key":"2858_CR32","doi-asserted-by":"publisher","unstructured":"Hellekes, J., M\u00fchlhaus, M., Bahmanyar, R., Azimi, S. M., & Kurz, F. (2024). VETRA: A Dataset for Vehicle Tracking in Aerial Imagery \u2013 New Challenges for Multi-Object Tracking. In: Computer Vision \u2013 ECCV 2024: 18th European Conference, Milan, Italy, September 29\u2013October 4, 2024, Proceedings, Part LXXXV Berlin, Heidelberg: Springer-Verlag. p. 52\u201370. https:\/\/doi.org\/10.1007\/978-3-031-73013-9_4.","DOI":"10.1007\/978-3-031-73013-9_4"},{"key":"2858_CR33","unstructured":"Hornakova, A., Henschel, R., Rosenhahn, B., & Swoboda, P. (2020). Lifted disjoint paths with application in multiple object tracking. In: ICML PMLR; p. 4364\u20134375."},{"key":"2858_CR34","doi-asserted-by":"crossref","unstructured":"Hornakova, A., Kaiser, T., Swoboda, P., Rolinek, M., Rosenhahn, B., & Henschel, R. (2021). Making Higher Order MOT Scalable: An Efficient Approximate Solver for Lifted Disjoint Paths. In: ICCV. p. 6330\u20136340.","DOI":"10.1109\/ICCV48922.2021.00627"},{"key":"2858_CR35","doi-asserted-by":"crossref","unstructured":"Huang, H. W., Yang, C. Y., Sun, J., Kim, P. K., Kim, K. J., Lee, K., Huang, C.-I., & Hwang, J.-N. (2024). Iterative Scale-Up ExpansionIoU and Deep Features Association for Multi-Object Tracking in Sports. In: Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision. p. 163\u2013172.","DOI":"10.1109\/WACVW60836.2024.00024"},{"key":"2858_CR36","unstructured":"Huang, S., Hou, Y., Liu, L., Yu, X., & Shen, X. (2025). Real-Time Object Detection Meets DINOv3. arXiv."},{"issue":"2","key":"2858_CR37","doi-asserted-by":"publisher","first-page":"319","DOI":"10.1109\/TPAMI.2008.57","volume":"31","author":"R Kasturi","year":"2009","unstructured":"Kasturi, R., Goldgof, D., Soundararajan, P., Manohar, V., Garofolo, J., Bowers, R., Boonstra, M., Korzhova, V., & Zhang, J. (2009). Framework for Performance Evaluation of Face, Text, and Vehicle Detection and Tracking in Video: Data, Metrics, and Protocol. IEEE Transactions on Pattern Analysis and Machine Intelligence., 31(2), 319\u2013336. https:\/\/doi.org\/10.1109\/TPAMI.2008.57","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence."},{"key":"2858_CR38","doi-asserted-by":"publisher","unstructured":"Keval, H., & Sasse, M. A. (2008). to catch a thief \u2013 you need at least 8 frames per second: the impact of frame rates on user performance in a CCTV detection task. In: Proceedings of the 16th ACM International Conference on Multimedia MM \u201908, New York, NY, USA: Association for Computing Machinery; p. 941\u201394. https:\/\/doi.org\/10.1145\/1459359.1459527.","DOI":"10.1145\/1459359.1459527"},{"key":"2858_CR39","unstructured":"Kingma, D. P., & Ba, J. (2014). Adam: A Method for Stochastic Optimization. arXiv:1412.6980. https:\/\/api.semanticscholar.org\/CorpusID:6628106."},{"issue":"1\u20132","key":"2858_CR40","doi-asserted-by":"publisher","first-page":"83","DOI":"10.1002\/nav.3800020109","volume":"2","author":"HW Kuhn","year":"1955","unstructured":"Kuhn, H. W. (1955). The Hungarian Method for the Assignment Problem. Naval Research Logistics Quarterly., 2(1\u20132), 83\u201397. https:\/\/doi.org\/10.1002\/nav.3800020109","journal-title":"Naval Research Logistics Quarterly."},{"key":"2858_CR41","doi-asserted-by":"crossref","unstructured":"Leal-Taix\u00e9, L., Canton-Ferrer, C., & Schindler, K. (2016). Learning by tracking: Siamese CNN for robust target association. arXiv:1604.07866","DOI":"10.1109\/CVPRW.2016.59"},{"key":"2858_CR42","doi-asserted-by":"crossref","unstructured":"Leal-Taixe, L., Canton-Ferrer, C., & Schindler, K. (2016). Learning by Tracking: Siamese CNN for Robust Target Association. In: CVPRW.","DOI":"10.1109\/CVPRW.2016.59"},{"key":"2858_CR43","doi-asserted-by":"crossref","unstructured":"Leal-Taix\u00e9, L., Fenzi, M., Kuznetsova, A., Rosenhahn, B., & Savarese, S. (2014). Learning an Image-Based Motion Context for Multiple People Tracking. In: 2014 IEEE Conference on Computer Vision and Pattern Recognition. p. 3542\u20133549.","DOI":"10.1109\/CVPR.2014.453"},{"key":"2858_CR44","doi-asserted-by":"crossref","unstructured":"Li, J., Gao, X., & Jiang, T. (2020). Graph Networks for Multiple Object Tracking. In: Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision (WACV).","DOI":"10.1109\/WACV45572.2020.9093347"},{"key":"2858_CR45","doi-asserted-by":"crossref","unstructured":"Li, J., Gao, X., & Jiang, T. (2020). Graph Networks for Multiple Object Tracking. In: 2020 IEEE Winter Conference on Applications of Computer Vision (WACV). p. 708\u2013717.","DOI":"10.1109\/WACV45572.2020.9093347"},{"issue":"2","key":"2858_CR46","doi-asserted-by":"publisher","first-page":"751","DOI":"10.1145\/3296957.3173191","volume":"53","author":"SC Lin","year":"2018","unstructured":"Lin, S. C., Zhang, Y., Hsu, C. H., Skach, M., Haque, M. E., Tang, L., & Mars, J. (2018). The Architectural Implications of Autonomous Driving: Constraints and Acceleration. SIGPLAN Not., 53(2), 751\u201376. https:\/\/doi.org\/10.1145\/3296957.3173191","journal-title":"SIGPLAN Not."},{"issue":"2","key":"2858_CR47","doi-asserted-by":"publisher","first-page":"318","DOI":"10.1109\/TPAMI.2018.2858826","volume":"42","author":"TY Lin","year":"2020","unstructured":"Lin, T. Y., Goyal, P., Girshick, R., He, K., & Doll\u00e1r, P. (2020). Focal Loss for Dense Object Detection. IEEE Transactions on Pattern Analysis and Machine Intelligence., 42(2), 318\u2013327. https:\/\/doi.org\/10.1109\/TPAMI.2018.2858826","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence."},{"key":"2858_CR48","doi-asserted-by":"crossref","unstructured":"Liu, Q., Chu, Q., Liu, B., & Yu, N. (2020). GSM: Graph Similarity Model for Multi-Object Tracking. In: IJCAI. p. 530\u2013536.","DOI":"10.24963\/ijcai.2020\/74"},{"issue":"2","key":"2858_CR49","doi-asserted-by":"publisher","first-page":"548","DOI":"10.1007\/s11263-020-01375-2","volume":"129","author":"J Luiten","year":"2021","unstructured":"Luiten, J., Osep, A., Dendorfer, P., Torr, P., Geiger, A., Leal-Taix\u00e9, L., & Leibe, B. (2021). HOTA: A Higher Order Metric for Evaluating Multi-object Tracking. Int J Comput Vision., 129(2), 548\u201357. https:\/\/doi.org\/10.1007\/s11263-020-01375-2","journal-title":"Int J Comput Vision."},{"key":"2858_CR50","doi-asserted-by":"crossref","unstructured":"Luo, C., Yang, X., & Yuille, A. (2021). Exploring Simple 3D Multi-Object Tracking for Autonomous Driving. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV). p. 10488\u201310497.","DOI":"10.1109\/ICCV48922.2021.01032"},{"key":"2858_CR51","doi-asserted-by":"crossref","unstructured":"Lv, W., Huang, Y., Zhang, N., Lin, R. S., Han, M., & Zeng, D. (2024). DiffMOT: A Real-time Diffusion-based Multiple Object Tracker with Non-linear Prediction. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. p. 19321\u201319330.","DOI":"10.1109\/CVPR52733.2024.01828"},{"key":"2858_CR52","doi-asserted-by":"crossref","unstructured":"Maggiolino, G., Ahmad, A., Cao, J., & Kitani, K. (2023). Deep OC-Sort: Multi-Pedestrian Tracking by Adaptive Re-Identification. In: 2023 IEEE International Conference on Image Processing (ICIP). p. 3025\u20133029.","DOI":"10.1109\/ICIP49359.2023.10222576"},{"key":"2858_CR53","doi-asserted-by":"crossref","unstructured":". Meinhardt, T., Kirillov, A., Leal-Taixe, L.,& Feichtenhofer, C. (2022) TrackFormer: Multi-Object Tracking with Transformers. In: The IEEE Conference on Computer Vision and Pattern Recognition (CVPR).","DOI":"10.1109\/CVPR52688.2022.00864"},{"key":"2858_CR54","doi-asserted-by":"crossref","unstructured":"Milan, A., Rezatofighi, S. H., Dick, A., Reid, I., & Schindler, K. (2017). Online Multi-Target Tracking Using Recurrent Neural Networks. In: Proceedings of the Thirty-First AAAI Conference on Artificial Intelligence.","DOI":"10.1609\/aaai.v31i1.11194"},{"key":"2858_CR55","doi-asserted-by":"crossref","unstructured":"Pang, J., Qiu, L., Li, X., Chen, H., Li, Q., Darrell, T., & Yu, F. (2021). Quasi-Dense Similarity Learning for Multiple Object Tracking. In: IEEE\/CVF Conference on Computer Vision and Pattern Recognition.","DOI":"10.1109\/CVPR46437.2021.00023"},{"key":"2858_CR56","doi-asserted-by":"crossref","unstructured":"Qin, Z., Zhou, S., Wang, L., Duan, J., Hua, G., & Tang, W. (2023). MotionTrack: Learning Robust Short-term and Long-term Motions for Multi-Object Tracking. arXiv:2303.10404.","DOI":"10.1109\/CVPR52729.2023.01720"},{"key":"2858_CR57","doi-asserted-by":"crossref","unstructured":"Ren, H., Han, S., Ding, H., Zhang, Z., Wang, H., & Wang, F. (2023). Focus On Details: Online Multi-Object Tracking with Diverse Fine-Grained Representation. In: 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). p. 11289\u201311298.","DOI":"10.1109\/CVPR52729.2023.01086"},{"key":"2858_CR58","doi-asserted-by":"crossref","unstructured":"Ristani, E., Solera, F., Zou R. S., Cucchiara, R., & Tomasi, C. (2016). Performance Measures and a Data Set for Multi-Target, Multi-Camera Tracking. arXiv:1609.01775","DOI":"10.1007\/978-3-319-48881-3_2"},{"key":"2858_CR59","doi-asserted-by":"crossref","unstructured":"Ristani, E., & Tomasi, C. (2018). Features for Multi-Target Multi-Camera Tracking and Re-Identification. In: CVPR.","DOI":"10.1109\/CVPR.2018.00632"},{"key":"2858_CR60","unstructured":"Robinson, I., Robicheaux, P., Popov, M., Ramanan, D., & Peri, N. (2025). RF-DETR: Neural Architecture Search for Real-Time Detection Transformers. arXiv:2511.09554."},{"key":"2858_CR61","doi-asserted-by":"crossref","unstructured":"Sadeghian, A., Alahi, A., & Savarese, S. (2017). Tracking the Untrackable: Learning to Track Multiple Cues With Long-Term Dependencies. In: ICCV.","DOI":"10.1109\/ICCV.2017.41"},{"key":"2858_CR62","unstructured":"Schmidt, F. (2012). Data Set for Tracking Vehicles in Aerial Image Sequences. KIT - Institute of Photogrammetry and Remote Sensing (IPF). https:\/\/www.ipf.kit.edu\/downloads_data_set_AIS_vehicle_tracking.php."},{"key":"2858_CR63","doi-asserted-by":"crossref","unstructured":"Seidenschwarz, J., Bras\u00f3, G., Serrano, V. C., Elezi, I., & Leal-Taix\u00e9, L. (2023). Simple Cues Lead to a Strong Multi-Object Tracker. arXiv:2206.04656.","DOI":"10.1109\/CVPR52729.2023.01327"},{"key":"2858_CR64","unstructured":"Shao, S., Zhao, Z., Li, B., Xiao, T., Yu, G., Zhang, X., & Sun, J. (2018). CrowdHuman: A Benchmark for Detecting Human in a Crowd. arXiv:1805.00123."},{"key":"2858_CR65","doi-asserted-by":"publisher","unstructured":"Somers, V., Alahi, A., & Vleeschouwer, C. D. (2024). Keypoint Promptable Re-Identification. In: Computer Vision - ECCV 2024 - 18th European Conference, Milan, Italy, September 29-October 4, 2024, Proceedings, Part LXXIX, vol. 15137 of Lecture Notes in Computer Science Springer. p. 216\u2013233. https:\/\/doi.org\/10.1007\/978-3-031-72986-7_13.","DOI":"10.1007\/978-3-031-72986-7_13"},{"key":"2858_CR66","doi-asserted-by":"crossref","unstructured":"Son, J., Baek, M., Cho, M., & Han, B. (2017). Multi-Object Tracking With Quadruplet Convolutional Neural Networks. In: CVPR.","DOI":"10.1109\/CVPR.2017.403"},{"key":"2858_CR67","doi-asserted-by":"crossref","unstructured":"Sun, P., Cao, J., Jiang, Y., Yuan, Z., Bai, S., Kitani, K., & Luo, P. (2022). DanceTrack: Multi-Object Tracking in Uniform Appearance and Diverse Motion. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR).","DOI":"10.1109\/CVPR52688.2022.02032"},{"key":"2858_CR68","doi-asserted-by":"crossref","unstructured":"Takala, V., & Pietikainen, M. (2007). Multi-Object Tracking Using Color, Texture and Motion. In: 2007 IEEE Conference on Computer Vision and Pattern Recognition. p. 1\u20137.","DOI":"10.1109\/CVPR.2007.383506"},{"key":"2858_CR69","doi-asserted-by":"crossref","unstructured":"Tang, S., Andres, B., Andriluka, M., & Schiele, B. (2015). Subgraph Decomposition for Multi-Target Tracking. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR).","DOI":"10.1109\/CVPR.2015.7299138"},{"key":"2858_CR70","doi-asserted-by":"crossref","unstructured":"Tang, S., Andriluka, M., Andres, B., & Schiele, B. (2017). Multiple People Tracking by Lifted Multicut and Person Re-Identification. In: CVPR.","DOI":"10.1109\/CVPR.2017.394"},{"key":"2858_CR71","doi-asserted-by":"publisher","first-page":"610","DOI":"10.1007\/978-3-642-15561-1_44","volume-title":"Computer Vision - ECCV 2010 Berlin","author":"C Vondrick","year":"2010","unstructured":"Vondrick, C., Ramanan, D., & Patterson, D. (2010). Efficiently Scaling Up Video Annotation with Crowdsourced Marketplaces. In K. Daniilidis, P. Maragos, & N. Paragios (Eds.), Computer Vision - ECCV 2010 Berlin (pp. 610\u2013623). Heidelberg: Springer, Berlin Heidelberg."},{"key":"2858_CR72","doi-asserted-by":"crossref","unstructured":"Wang, G., Yang, S., Liu, H., Wang, Z., Yang, Y., Wang, S., Yu, G., Zhou, E., & Sun, J. (2020). High-Order Information Matters: Learning Relation and Topology for Occluded Person Re-Identification. arXiv:2003.08177.","DOI":"10.1109\/CVPR42600.2020.00648"},{"key":"2858_CR73","doi-asserted-by":"publisher","unstructured":"Wang, Y., Kitani, K., & Weng, X. (2021). Joint Object Detection and Multi-Object Tracking with Graph Neural Networks. In: 2021 IEEE International Conference on Robotics and Automation (ICRA) IEEE Press. p. 13708\u201313715. https:\/\/doi.org\/10.1109\/ICRA48506.2021.9561110.","DOI":"10.1109\/ICRA48506.2021.9561110"},{"key":"2858_CR74","doi-asserted-by":"crossref","unstructured":"Weng, X., Wang, Y., Man, Y., & Kitani, K. M. (2020). GNN3DMOT: Graph Neural Network for 3D Multi-Object Tracking With 2D-3D Multi-Feature Learning. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR).","DOI":"10.1109\/CVPR42600.2020.00653"},{"key":"2858_CR75","doi-asserted-by":"crossref","unstructured":"Wojke, N., Bewley, A., & Paulus, D. (2017). Simple Online and Realtime Tracking with a Deep Association Metric. In: 2017 IEEE International Conference on Image Processing (ICIP) IEEE. p. 3645\u20133649.","DOI":"10.1109\/ICIP.2017.8296962"},{"issue":"1","key":"2858_CR76","doi-asserted-by":"publisher","first-page":"275","DOI":"10.1109\/TCSVT.2020.2975842","volume":"31","author":"J Xiang","year":"2021","unstructured":"Xiang, J., Xu, G., Ma, C., & Hou, J. (2021). End-to-End Learning Deep CRF Models for Multi-Object Tracking Deep CRF Models. IEEE Transactions on Circuits and Systems for Video Technology., 31(1), 275\u2013288. https:\/\/doi.org\/10.1109\/TCSVT.2020.2975842","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology."},{"key":"2858_CR77","doi-asserted-by":"publisher","unstructured":"Xiao, C., Cao, Q., Luo, Z., & Lan ,L. (2024). MambaTrack: A Simple Baseline for Multiple Object Tracking with State Space Model. In: Proceedings of the 32nd ACM International Conference on Multimedia MM \u201924, New York, NY, USA: Association for Computing Machinery. p. 4082\u2013409. https:\/\/doi.org\/10.1145\/3664647.3680944.","DOI":"10.1145\/3664647.3680944"},{"key":"2858_CR78","doi-asserted-by":"publisher","first-page":"28","DOI":"10.1007\/s00464-014-3504-z","volume":"03","author":"S Xu","year":"2014","unstructured":"Xu, S., Perez, M., Yang, K., Perrenot, C., Felblinger, J., & Hubert, J. (2014). Determination of the latency effects on surgical performance and the acceptable latency levels in telesurgery using the dV-Trainer (R) simulator. Surgical endoscopy., 03, 28. https:\/\/doi.org\/10.1007\/s00464-014-3504-z","journal-title":"Surgical endoscopy."},{"key":"2858_CR79","doi-asserted-by":"crossref","unstructured":"Xu, Y., Osep, A., Ban, Y., Horaud, R., Leal-Taix\u00e9, L., & Alameda-Pineda, X. (2020). How To Train Your Deep Multi-Object Tracker. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition. p. 6787\u20136796.","DOI":"10.1109\/CVPR42600.2020.00682"},{"key":"2858_CR80","doi-asserted-by":"crossref","unstructured":"Yang, M., Han, G., Yan, B., Zhang, W., Qi, J., Lu, H., & Wang, D. (2024). Hybrid-sort: Weak cues matter for online multi-object tracking. In: Proceedings of the AAAI Conference on Artificial Intelligence, (Vol.\u00a038, pp. 6504\u20136512).","DOI":"10.1609\/aaai.v38i7.28471"},{"key":"2858_CR81","doi-asserted-by":"crossref","unstructured":"You, S., Yao, H., Bao, B. k., & Xu, C. (2023). UTM: A Unified Multiple Object Tracking Model with Identity-Aware Feature Enhancement. In: 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). p. 21876\u201321886.","DOI":"10.1109\/CVPR52729.2023.02095"},{"key":"2858_CR82","doi-asserted-by":"crossref","unstructured":"Yu, F., Chen, H., Wang, X., Xian, W., Chen, Y., Liu, F., Madhavan, V., & Darrell, T. (2020). BDD100K: A Diverse Driving Dataset for Heterogeneous Multitask Learning. In: IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR).","DOI":"10.1109\/CVPR42600.2020.00271"},{"key":"2858_CR83","doi-asserted-by":"crossref","unstructured":"Zamir, A. R., Dehghan, A., & Shah, M. (2012). Gmcp-tracker: Global multi-object tracking using generalized minimum clique graphs. In: ECCV Springer. p. 343\u2013356.","DOI":"10.1007\/978-3-642-33709-3_25"},{"key":"2858_CR84","doi-asserted-by":"publisher","unstructured":"Zeng, F., Dong, B., Zhang, Y., Wang, T., Zhang, X., & Wei, Y. (2022). MOTR: End-to-End Multiple-Object Tracking with Transformer. In: Computer Vision \u2013 ECCV 2022: 17th European Conference, Tel Aviv, Israel, October 23\u201327, 2022, Proceedings, Part XXVII Berlin, Heidelberg: Springer-Verlag. p. 659\u201367. https:\/\/doi.org\/10.1007\/978-3-031-19812-0_38.","DOI":"10.1007\/978-3-031-19812-0_38"},{"key":"2858_CR85","unstructured":"Zhang, L., Li, Y., & Nevatia, R. (2008). Global data association for multi-object tracking using network flows. In: CVPR."},{"key":"2858_CR86","doi-asserted-by":"crossref","unstructured":"Zhang, L., Li, Y., & Nevatia, R. (2008). Global data association for multi-object tracking using network flows. In: 2008 IEEE Conference on Computer Vision and Pattern Recognition IEEE. p. 1\u20138.","DOI":"10.1109\/CVPR.2008.4587584"},{"key":"2858_CR87","doi-asserted-by":"crossref","unstructured":"Zhang, Y., Sun, P., Jiang, Y., Yu, D., Weng, F., Yuan, Z., Luo, P., Liu, W., & Wang, X. (2022). ByteTrack: Multi-Object Tracking by Associating Every Detection Box. Proceedings of the European Conference on Computer Vision (ECCV).","DOI":"10.1007\/978-3-031-20047-2_1"},{"key":"2858_CR88","doi-asserted-by":"publisher","first-page":"3069","DOI":"10.1007\/s11263-021-01513-4","volume":"129","author":"Y Zhang","year":"2021","unstructured":"Zhang, Y., Wang, C., Wang, X., Zeng, W., & Liu, W. (2021). Fairmot: On the fairness of detection and re-identification in multiple object tracking. International Journal of Computer Vision., 129, 3069\u20133087.","journal-title":"International Journal of Computer Vision."},{"key":"2858_CR89","doi-asserted-by":"publisher","unstructured":"Zhang, Y., Wang, T., & Zhang. X. (2023). MOTRv2: Bootstrapping End-to-End Multi-Object Tracking by Pretrained Object Detectors. In: 2023 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR) IEEE. p. 22056\u201322065.https:\/\/doi.org\/10.1109\/CVPR52729.2023.02112.","DOI":"10.1109\/CVPR52729.2023.02112"},{"key":"2858_CR90","doi-asserted-by":"crossref","unstructured":"Zhou, X., Koltun, V., & Kr\u00e4henb\u00fchl, P. (2020). Tracking Objects as Points. ECCV.","DOI":"10.1007\/978-3-030-58548-8_28"},{"key":"2858_CR91","doi-asserted-by":"crossref","unstructured":"Zhou, X., Yin, T., Koltun, V., & Kr\u00e4henb\u00fchl, P. (2022). Global Tracking Transformers. In: CVPR.","DOI":"10.1109\/CVPR52688.2022.00857"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02858-4.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-026-02858-4","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02858-4.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T16:15:09Z","timestamp":1784564109000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-026-02858-4"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,15]]},"references-count":91,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2858"],"URL":"https:\/\/doi.org\/10.1007\/s11263-026-02858-4","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,15]]},"assertion":[{"value":"4 September 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"15 April 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"15 May 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}}],"article-number":"272"}}