{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,29]],"date-time":"2026-05-29T11:28:05Z","timestamp":1780054085971,"version":"3.54.0"},"reference-count":54,"publisher":"Tsinghua University Press","issue":"4","license":[{"start":{"date-parts":[[2023,12,1]],"date-time":"2023-12-01T00:00:00Z","timestamp":1701388800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0"},{"start":{"date-parts":[[2023,7,5]],"date-time":"2023-07-05T00:00:00Z","timestamp":1688515200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Comp. Visual. Med."],"published-print":{"date-parts":[[2023,12]]},"DOI":"10.1007\/s41095-022-0302-8","type":"journal-article","created":{"date-parts":[[2023,7,20]],"date-time":"2023-07-20T00:01:42Z","timestamp":1689811302000},"page":"753-765","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":10,"title":["Sequential interactive image segmentation"],"prefix":"10.26599","volume":"9","author":[{"given":"Zheng","family":"Lin","sequence":"first","affiliation":[{"name":"TKLNDST, College of Computer Science, Nankai University, Tianjin, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zhao","family":"Zhang","sequence":"additional","affiliation":[{"name":"TKLNDST, College of Computer Science, Nankai University, Tianjin, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zi-Yue","family":"Zhu","sequence":"additional","affiliation":[{"name":"TKLNDST, College of Computer Science, Nankai University, Tianjin, China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Deng-Ping","family":"Fan","sequence":"additional","affiliation":[{"name":"Computer Vision Lab, ETH Zurich, Zurich, Switzerland"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xia-Lei","family":"Liu","sequence":"additional","affiliation":[{"name":"TKLNDST, College of Computer Science, Nankai University, Tianjin, China"}],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"11138","reference":[{"key":"302_CR1","doi-asserted-by":"crossref","unstructured":"Maninis, K. K.; Caelles, S.; Pont-Tuset, J.; Van Gool, L. Deep extreme cut: From extreme points to object segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 616\u2013625, 2018.","DOI":"10.1109\/CVPR.2018.00071"},{"key":"302_CR2","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"20","DOI":"10.1007\/978-3-030-01264-9_2","volume-title":"Computer Vision\u2013ECCV 2018","author":"H Le","year":"2018","unstructured":"Le, H.; Mai, L.; Price, B.; Cohen, S.; Jin, H. L.; Liu, F. Interactive boundary prediction for object selection. In: Computer Vision\u2013ECCV 2018. Lecture Notes in Computer Science, Vol. 11218. Ferrari, V.; Hebert, M.; Sminchisescu, C.; Weiss, Y. Eds. Springer Cham, 20\u201336, 2018."},{"issue":"9","key":"302_CR3","doi-asserted-by":"publisher","first-page":"1321","DOI":"10.1007\/s11263-019-01184-2","volume":"127","author":"S D Jain","year":"2019","unstructured":"Jain, S. D.; Grauman, K. Click carving: Interactive object segmentation in images and videos with point clicks. International Journal of Computer Vision Vol. 127, No. 9, 1321\u20131344, 2019.","journal-title":"International Journal of Computer Vision"},{"key":"302_CR4","doi-asserted-by":"crossref","unstructured":"Xu, N.; Price, B.; Cohen, S.; Yang, J. M.; Huang, T. Deep GrabCut for object selection. In: Proceedings of the British Machine Vision Conference, 182.1\u2013182.12, 2017.","DOI":"10.5244\/C.31.182"},{"key":"302_CR5","unstructured":"Majumder, S.; Rai, A.; Khurana, A.; Yao, A. Two-in-one refinement for interactive segmentation. In: Proceedings of the 31st British Machine Vision Conference, 2020."},{"key":"302_CR6","doi-asserted-by":"crossref","unstructured":"Zhang, S. Y.; Liew, J. H.; Wei, Y. C.; Wei, S. K.; Zhao, Y. Interactive object segmentation with inside-outside guidance. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 12231\u201312241, 2020.","DOI":"10.1109\/CVPR42600.2020.01225"},{"key":"302_CR7","doi-asserted-by":"crossref","unstructured":"Li, Z. W.; Chen, Q. F.; Koltun, V. Interactive image segmentation with latent diversity. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 577\u2013585, 2018.","DOI":"10.1109\/CVPR.2018.00067"},{"key":"302_CR8","doi-asserted-by":"crossref","unstructured":"Liew, J. H.; Cohen, S.; Price, B.; Mai, L.; Ong, S. H.; Feng, J. S. MultiSeg: Semantically meaningful, scale-diverse segmentations from minimal user input. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 662\u2013670, 2019.","DOI":"10.1109\/ICCV.2019.00075"},{"key":"302_CR9","unstructured":"Mahadevan, S.; Voigtlaender, P.; Leibe, B. Iteratively trained interactive segmentation. arXiv preprint arXiv:1805.04398, 2018."},{"key":"302_CR10","doi-asserted-by":"crossref","unstructured":"Majumder, S.; Yao, A. Content-aware multi-level guidance for interactive instance segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 11594\u201311603, 2019.","DOI":"10.1109\/CVPR.2019.01187"},{"key":"302_CR11","doi-asserted-by":"crossref","unstructured":"Lin, Z.; Zhang, Z.; Chen, L. Z.; Cheng, M. M.; Lu, S. P. Interactive image segmentation with first click attention. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 13336\u201313345, 2020.","DOI":"10.1109\/CVPR42600.2020.01335"},{"key":"302_CR12","doi-asserted-by":"crossref","unstructured":"Jang, W. D.; Kim, C. S. Interactive image segmentation via backpropagating refinement scheme. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 5292\u20135301, 2019.","DOI":"10.1109\/CVPR.2019.00544"},{"key":"302_CR13","doi-asserted-by":"crossref","unstructured":"Sofiiuk, K.; Petrov, I.; Barinova, O.; Konushin, A. F-BRS: Rethinking backpropagating refinement for interactive segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 8620\u20138629, 2020.","DOI":"10.1109\/CVPR42600.2020.00865"},{"key":"302_CR14","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"579","DOI":"10.1007\/978-3-030-58517-4_34","volume-title":"Computer Vision\u2013ECCV 2020","author":"T Kontogianni","year":"2020","unstructured":"Kontogianni, T.; Gygli, M.; Uijlings, J.; Ferrari, V. Continuous adaptation for interactive object segmentation by learning from corrections. In: Computer Vision\u2013ECCV 2020. Lecture Notes in Computer Science, Vol. 12361. Vedaldi, A.; Bischof, H.; Brox, T.; Frahm, J. M. Eds. Springer Cham, 579\u2013596, 2020."},{"issue":"1","key":"302_CR15","doi-asserted-by":"publisher","first-page":"165","DOI":"10.1007\/s41095-021-0235-7","volume":"8","author":"L X Gong","year":"2022","unstructured":"Gong, L. X.; Zhang, Y. Q.; Zhang, Y. K.; Yang, Y.; Xu, W. W. Erroneous pixel prediction for semantic image segmentation. Computational Visual Media Vol. 8, No. 1, 165\u2013175, 2022.","journal-title":"Computational Visual Media"},{"issue":"11","key":"302_CR16","doi-asserted-by":"publisher","first-page":"219101","DOI":"10.1007\/s11432-019-2759-y","volume":"63","author":"X Y Zhang","year":"2020","unstructured":"Zhang, X. Y.; Wang, L. J.; Xie, J.; Zhu, P. F. Human-in-the-loop image segmentation and annotation. Science China Information Sciences Vol. 63, No. 11, 219101, 2020.","journal-title":"Science China Information Sciences"},{"issue":"4","key":"302_CR17","first-page":"150","volume":"1","author":"V Vezhnevets","year":"2005","unstructured":"Vezhnevets, V.; Konouchine, V. \u201cGrowCut\u201d - Interactive multi-label N-D image segmentation by cellular automata. Proc. of Graph. Vol. 1, No. 4, 150\u2013156, 2005.","journal-title":"Proc. of Graph."},{"issue":"2","key":"302_CR18","doi-asserted-by":"publisher","first-page":"113","DOI":"10.1007\/s11263-008-0191-z","volume":"82","author":"X Bai","year":"2009","unstructured":"Bai, X.; Sapiro, G. Geodesic matting: A framework for fast interactive image and video segmentation and matting. International Journal of Computer Vision Vol. 82, No. 2, 113\u2013132, 2009.","journal-title":"International Journal of Computer Vision"},{"key":"302_CR19","doi-asserted-by":"crossref","unstructured":"Gulshan, V.; Rother, C.; Criminisi, A.; Blake, A.; Zisserman, A. Geodesic star convexity for interactive image segmentation. In: Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition, 3129\u20133136, 2010.","DOI":"10.1109\/CVPR.2010.5540073"},{"key":"302_CR20","doi-asserted-by":"crossref","unstructured":"Kim, T. H.; Lee, K. M.; Lee, S. U. Nonparametric higher-order learning for interactive segmentation. In: Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition, 3201\u20133208, 2010.","DOI":"10.1109\/CVPR.2010.5540078"},{"issue":"3","key":"302_CR21","doi-asserted-by":"publisher","first-page":"1301","DOI":"10.1109\/TIP.2016.2518480","volume":"25","author":"M Jian","year":"2016","unstructured":"Jian, M.; Jung, C. Interactive image segmentation using adaptive constraint propagation. IEEE Transactions on Image Processing Vol. 25, No. 3, 1301\u20131311, 2016.","journal-title":"IEEE Transactions on Image Processing"},{"issue":"1","key":"302_CR22","doi-asserted-by":"publisher","first-page":"330","DOI":"10.1109\/TIP.2018.2867941","volume":"28","author":"T Wang","year":"2019","unstructured":"Wang, T.; Yang, J.; Ji, Z. X.; Sun, Q. S. Probabilistic diffusion for interactive image segmentation. IEEE Transactions on Image Processing Vol. 28, No. 1, 330\u2013342, 2019.","journal-title":"IEEE Transactions on Image Processing"},{"key":"302_CR23","doi-asserted-by":"crossref","unstructured":"Wu, J. J.; Zhao, Y. B.; Zhu, J. Y.; Luo, S. W.; Tu, Z. W. MILCut: A sweeping line multiple instance learning paradigm for interactive image segmentation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 256\u2013263, 2014.","DOI":"10.1109\/CVPR.2014.40"},{"key":"302_CR24","doi-asserted-by":"crossref","unstructured":"Bai, J. J.; Wu, X. D. Error-tolerant scribbles based interactive image segmentation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 392\u2013399, 2014.","DOI":"10.1109\/CVPR.2014.57"},{"key":"302_CR25","doi-asserted-by":"crossref","unstructured":"Rother, C.; Kolmogorov, V.; Blake, A. \u201cGrabCut\u201d: Interactive foreground extraction using iterated graph cuts. In: Proceedings of the ACM SIGGRAPH 2004 Papers, 309\u2013314, 2004.","DOI":"10.1145\/1186562.1015720"},{"key":"302_CR26","doi-asserted-by":"crossref","unstructured":"Mortensen, E. N.; Barrett, W. A. Intelligent scissors for image composition. In: Proceedings of the 22nd Annual Conference on Computer Graphics and Interactive Techniques, 191\u2013198, 1995.","DOI":"10.1145\/218380.218442"},{"issue":"3","key":"302_CR27","doi-asserted-by":"publisher","first-page":"303","DOI":"10.1145\/1015706.1015719","volume":"23","author":"Y Li","year":"2004","unstructured":"Li, Y.; Sun, J. A.; Tang, C. K.; Shum, H. Y. Lazy snapping. ACM Transactions on Graphics Vol. 23, No. 3, 303\u2013308, 2004.","journal-title":"ACM Transactions on Graphics"},{"key":"302_CR28","doi-asserted-by":"crossref","unstructured":"Xu, N.; Price, B.; Cohen, S.; Yang, J. M.; Huang, T. Deep interactive object selection. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 373\u2013381, 2016.","DOI":"10.1109\/CVPR.2016.47"},{"key":"302_CR29","doi-asserted-by":"crossref","unstructured":"Boykov, Y. Y.; Jolly, M.-P. Interactive graph cuts for optimal boundary & region segmentation of objects in N-D images. In: Proceedings of the 8th IEEE International Conference on Computer Vision, 105\u2013112, 2001.","DOI":"10.1109\/ICCV.2001.937505"},{"issue":"9","key":"302_CR30","doi-asserted-by":"publisher","first-page":"1124","DOI":"10.1109\/TPAMI.2004.60","volume":"26","author":"Y Boykov","year":"2004","unstructured":"Boykov, Y.; Kolmogorov, V. An experimental comparison of min-cut\/max- flow algorithms for energy minimization in vision. IEEE Transactions on Pattern Analysis and Machine Intelligence Vol. 26, No. 9, 1124\u20131137, 2004.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"issue":"11","key":"302_CR31","doi-asserted-by":"publisher","first-page":"1768","DOI":"10.1109\/TPAMI.2006.233","volume":"28","author":"L Grady","year":"2006","unstructured":"Grady, L. Random walks for image segmentation. IEEE Transactions on Pattern Analysis and Machine Intelligence Vol. 28, No. 11, 1768\u20131783, 2006.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"302_CR32","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"264","DOI":"10.1007\/978-3-540-88690-7_20","volume-title":"Computer Vision\u2013ECCV 2008","author":"T H Kim","year":"2008","unstructured":"Kim, T. H.; Lee, K. M.; Lee, S. U. Generative image segmentation using random walks with restart. In: Computer Vision\u2013ECCV 2008. Lecture Notes in Computer Science, Vol. 5304. Forsyth, D.; Torr, P.; Zisserman, A. Eds. Springer Berlin Heidelberg, 264\u2013275, 2008."},{"key":"302_CR33","doi-asserted-by":"crossref","unstructured":"Castrej\u00f3n, L.; Kundu, K.; Urtasun, R.; Fidler, S. Annotating object instances with a polygon-RNN. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 4485\u20134493, 2017.","DOI":"10.1109\/CVPR.2017.477"},{"key":"302_CR34","doi-asserted-by":"crossref","unstructured":"Acuna, D.; Ling, H.; Kar, A.; Fidler, S. Efficient interactive annotation of segmentation datasets with polygon-RNN. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 859\u2013868, 2018.","DOI":"10.1109\/CVPR.2018.00096"},{"key":"302_CR35","doi-asserted-by":"crossref","unstructured":"Ling, H.; Gao, J.; Kar, A.; Chen, W. Z.; Fidler, S. Fast interactive object annotation with curve-GCN. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 5252\u20135261, 2019.","DOI":"10.1109\/CVPR.2019.00540"},{"key":"302_CR36","doi-asserted-by":"crossref","unstructured":"Lee, K. M.; Myeong, H.; Song, G. SeedNet: Automatic seed generation with deep reinforcement learning for robust interactive segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 1760\u20131768, 2018.","DOI":"10.1109\/CVPR.2018.00189"},{"key":"302_CR37","doi-asserted-by":"crossref","unstructured":"Liew, J.; Wei, Y. C.; Xiong, W.; Ong, S. H.; Feng, J. S. Regional interactive image segmentation networks. In: Proceedings of the IEEE International Conference on Computer Vision, 746\u20132754, 2017.","DOI":"10.1109\/ICCV.2017.297"},{"key":"302_CR38","doi-asserted-by":"publisher","first-page":"31","DOI":"10.1016\/j.neunet.2018.10.009","volume":"109","author":"Y Hu","year":"2019","unstructured":"Hu, Y.; Soltoggio, A.; Lock, R.; Carter, S. A fully convolutional two-stream fusion network for interactive image segmentation. Neural Networks Vol. 109, 31\u201342, 2019.","journal-title":"Neural Networks"},{"key":"302_CR39","doi-asserted-by":"crossref","unstructured":"Benenson, R.; Popov, S.; Ferrari, V. Large-scale interactive object segmentation with human annotators. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 11692\u201311701, 2019.","DOI":"10.1109\/CVPR.2019.01197"},{"key":"302_CR40","doi-asserted-by":"crossref","unstructured":"Lin, Z.; Duan, Z. P.; Zhang, Z.; Guo, C. L.; Cheng, M. M. FocusCut: Diving into a focus view in interactive segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 2627\u20132636, 2022.","DOI":"10.1109\/CVPR52688.2022.00266"},{"key":"302_CR41","doi-asserted-by":"crossref","unstructured":"Zhang, C. B.; Xiao, J. W.; Liu, X. L.; Chen, Y. C.; Cheng, M. M. Representation compensation networks for continual semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 7043\u20137054, 2022.","DOI":"10.1109\/CVPR52688.2022.00692"},{"key":"302_CR42","doi-asserted-by":"crossref","unstructured":"Cermelli, F.; Mancini, M.; Rota Bul\u00f2, S.; Ricci, E.; Caputo, B. Modeling the background for incremental learning in semantic segmentation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 9230\u20139239, 2020.","DOI":"10.1109\/CVPR42600.2020.00925"},{"key":"302_CR43","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"833","DOI":"10.1007\/978-3-030-01234-2_49","volume-title":"Computer Vision\u2013ECCV 2018","author":"L C Chen","year":"2018","unstructured":"Chen, L. C.; Zhu, Y. K.; Papandreou, G.; Schroff, F.; Adam, H. Encoder\u2013decoder with atrous separable convolution for semantic image segmentation. In: Computer Vision\u2013ECCV 2018. Lecture Notes in Computer Science, Vol. 11211. Ferrari, V.; Hebert, M.; Sminchisescu, C.; Weiss, Y. Eds. Springer Cham, 833\u2013851, 2018."},{"key":"302_CR44","doi-asserted-by":"crossref","unstructured":"He, K. M.; Zhang, X. Y.; Ren, S. Q.; Sun, J. Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 770\u2013778, 2016.","DOI":"10.1109\/CVPR.2016.90"},{"issue":"2","key":"302_CR45","doi-asserted-by":"publisher","first-page":"303","DOI":"10.1007\/s11263-009-0275-4","volume":"88","author":"M Everingham","year":"2010","unstructured":"Everingham, M.; Gool, L.; Williams, C. K. I.; Winn, J.; Zisserman, A. The pascal visual object classes (VOC) challenge. International Journal of Computer Vision Vol. 88, No. 2, 303\u2013338, 2010.","journal-title":"International Journal of Computer Vision"},{"key":"302_CR46","doi-asserted-by":"crossref","unstructured":"Hariharan, B.; Arbel\u00e1ez, P.; Bourdev, L.; Maji, S.; Malik, J. Semantic contours from inverse detectors. In: Proceedings of the International Conference on Computer Vision, 991\u2013998, 2011.","DOI":"10.1109\/ICCV.2011.6126343"},{"key":"302_CR47","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"740","DOI":"10.1007\/978-3-319-10602-1_48","volume-title":"Computer Vision\u2013ECCV 2014","author":"T Y Lin","year":"2014","unstructured":"Lin, T. Y.; Maire, M.; Belongie, S.; Hays, J.; Perona, P.; Ramanan, D.; Doll\u00e1r, P.; Zitnick, C. L. Microsoft COCO: Common objects in context. In: Computer Vision\u2013ECCV 2014. Lecture Notes in Computer Science, Vol. 8693. Fleet, D.; Pajdla, T.; Schiele, B.; Tuytelaars, T. Eds. Springer Cham, 740\u2013755, 2014."},{"key":"302_CR48","doi-asserted-by":"crossref","unstructured":"Fan, D. P.; Lin, Z.; Ji, G. P.; Zhang, D. W.; Fu, H. Z.; Cheng, M. M. Taking a deeper look at co-salient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 2916\u20132926, 2020.","DOI":"10.1109\/CVPR42600.2020.00299"},{"issue":"8","key":"302_CR49","first-page":"4339","volume":"44","author":"D P Fan","year":"2022","unstructured":"Fan, D. P.; Li, T. P.; Lin, Z.; Ji, G. P.; Zhang, D. W.; Cheng, M. M.; Fu, H. Z.; Shen, J. B. Re-thinking co-salient object detection. IEEE Transactions on Pattern Analysis and Machine Intelligence Vol. 44, No. 8, 4339\u20134354, 2022.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"302_CR50","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"crossref","first-page":"455","DOI":"10.1007\/978-3-030-58610-2_27","volume-title":"Computer Vision\u2013ECCV 2020","author":"Z Zhang","year":"2020","unstructured":"Zhang, Z.; Jin, W. D.; Xu, J.; Cheng, M. M. Gradient-induced co-saliency detection. In: Computer Vision\u2013ECCV 2020. Lecture Notes in Computer Science, Vol. 12357. Vedaldi, A.; Bischof, H.; Brox, T.; Frahm, J. M. Eds. Springer Cham, 455\u2013472, 2020."},{"key":"302_CR51","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"316","DOI":"10.1007\/978-3-030-58452-8_19","volume-title":"Computer Vision\u2013ECCV 2020","author":"M L Jia","year":"2020","unstructured":"Jia, M. L.; Shi, M. Y.; Sirotenko, M.; Cui, Y.; Cardie, C.; Hariharan, B.; Adam, H.; Belongie, S. Fashionpedia: Ontology, segmentation, and an attribute localization dataset. In: Computer Vision\u2013ECCV 2020. Lecture Notes in Computer Science, Vol. 12346. Vedaldi, A.; Bischof, H.; Brox, T.; Frahm, J. M. Eds. Springer Cham, 316\u2013332, 2020."},{"key":"302_CR52","doi-asserted-by":"crossref","unstructured":"Wang, J.; Markert, K.; Everingham, M. Learning models for object recognition from natural language descriptions. In: Proceedings of the British Machine Vision Conference, 2.1\u20132.11, 2009.","DOI":"10.5244\/C.23.2"},{"key":"302_CR53","doi-asserted-by":"crossref","unstructured":"Deng, J.; Dong, W.; Socher, R.; Li, L. J.; Kai, L.; Li, F. F. ImageNet: A large-scale hierarchical image database. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 248\u2013255, 2009.","DOI":"10.1109\/CVPR.2009.5206848"},{"key":"302_CR54","unstructured":"Steiner, B.; DeVito, Z.; Chintala, S.; Gross, S.; Paszke, A.; Massa, F.; Lerer, A.; Chanan, G.; Lin, Z.; Yang, E.; et al. PyTorch: An imperative style, high-performance deep learning library. In: Proceedings of the 33rd International Conference on Neural Information Processing Systems, Article No. 721, 8026\u20138037, 2019."}],"container-title":["Computational Visual Media"],"original-title":[],"link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s41095-022-0302-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s41095-022-0302-8\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s41095-022-0302-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"},{"URL":"http:\/\/xplorestaging.ieee.org\/ielx8\/10750449\/10897724\/10897731.pdf?arnumber=10897731","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,6,25]],"date-time":"2025-06-25T18:22:03Z","timestamp":1750875723000},"score":1,"resource":{"primary":{"URL":"https:\/\/ieeexplore.ieee.org\/document\/10897731\/"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,12]]},"references-count":54,"journal-issue":{"issue":"4"},"URL":"https:\/\/doi.org\/10.1007\/s41095-022-0302-8","relation":{},"ISSN":["2096-0662","2096-0433"],"issn-type":[{"value":"2096-0662","type":"electronic"},{"value":"2096-0433","type":"print"}],"subject":[],"published":{"date-parts":[[2023,12]]},"assertion":[{"value":"11 February 2022","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"17 June 2022","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"5 July 2023","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declaration of competing interest"}}]}}