{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,2,19]],"date-time":"2026-02-19T04:34:12Z","timestamp":1771475652593,"version":"3.50.1"},"reference-count":41,"publisher":"Springer Science and Business Media LLC","issue":"20","license":[{"start":{"date-parts":[[2024,8,2]],"date-time":"2024-08-02T00:00:00Z","timestamp":1722556800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,8,2]],"date-time":"2024-08-02T00:00:00Z","timestamp":1722556800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62006098"],"award-info":[{"award-number":["62006098"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62176106"],"award-info":[{"award-number":["62176106"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62076124"],"award-info":[{"award-number":["62076124"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"name":"The Key Projects of the National Natural Science Foundation of China","award":["U1836220"],"award-info":[{"award-number":["U1836220"]}]},{"name":"Fellowship of China Postdoctoral Science Foundation","award":["2020M681515"],"award-info":[{"award-number":["2020M681515"]}]},{"name":"Jiangsu Province Key Research and Development Plan","award":["BE2020036"],"award-info":[{"award-number":["BE2020036"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-024-19915-0","type":"journal-article","created":{"date-parts":[[2024,8,2]],"date-time":"2024-08-02T04:12:33Z","timestamp":1722571953000},"page":"22251-22268","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Label correlation preserving visual-semantic joint embedding for multi-label zero-shot learning"],"prefix":"10.1007","volume":"84","author":[{"given":"Zhongchen","family":"Ma","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Runze","family":"Ma","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Guangchen","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Qirong","family":"Mao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ming","family":"Dong","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,8,2]]},"reference":[{"key":"19915_CR1","doi-asserted-by":"crossref","unstructured":"Chen T, Pu T, Wu H, Xie Y, Lin L (2022) Structured semantic transfer for multi-label recognition with partial labels. In: Proceedings of the AAAI conference on artificial intelligence, vol 36, pp 339\u2013346","DOI":"10.1609\/aaai.v36i1.19910"},{"key":"19915_CR2","doi-asserted-by":"crossref","unstructured":"Coulibaly S, Kamsu-Foguem B, Kamissoko D, Traore D (2022) Deep convolution neural network sharing for the multi-label images classification. Mach Learn Appl 10:100422","DOI":"10.1016\/j.mlwa.2022.100422"},{"key":"19915_CR3","doi-asserted-by":"crossref","unstructured":"Zhao X, An Y, Xu N, Geng X (2022) Fusion label enhancement for multi-label learning. In: Proceedings of the thirty-first international joint conference on artificial intelligence, IJCAI","DOI":"10.24963\/ijcai.2022\/524"},{"key":"19915_CR4","doi-asserted-by":"crossref","unstructured":"Verelst T, Rubenstein PK, Eichner M, Tuytelaars T, Berman M (2023) Spatial consistency loss for training multi-label classifiers from single-label annotations. In: Proceedings of the IEEE\/CVF winter conference on applications of computer vision, pp 3879\u20133889","DOI":"10.1109\/WACV56688.2023.00387"},{"issue":"3","key":"19915_CR5","doi-asserted-by":"publisher","first-page":"697","DOI":"10.1007\/s13042-022-01658-9","volume":"14","author":"M Han","year":"2023","unstructured":"Han M, Wu H, Chen Z, Li M, Zhang X (2023) A survey of multi-label classification based on supervised and semi-supervised learning. Int J Mach Learn Cybern 14(3):697\u2013724","journal-title":"Int J Mach Learn Cybern"},{"key":"19915_CR6","doi-asserted-by":"crossref","unstructured":"Xie G-S, Liu L, Zhu F, Zhao F, Zhang Z, Yao Y, Qin J, Shao L (2020) Region graph embedding network for zero-shot learning. In: European conference on computer vision, Springer, pp 562\u2013580","DOI":"10.1007\/978-3-030-58548-8_33"},{"key":"19915_CR7","doi-asserted-by":"crossref","unstructured":"Chen Z, Huang Y, Chen J, Geng Y, Zhang W, Fang Y, Pan JZ, Song W, Chen H (2022) Duet: Cross-modal semantic grounding for contrastive zero-shot learning. arXiv preprint arXiv:2207.01328","DOI":"10.1609\/aaai.v37i1.25114"},{"key":"19915_CR8","doi-asserted-by":"crossref","unstructured":"Guo J, Guo S, Zhou Q, Liu Z, Lu X, Huo F (2023) Graph knows unknowns: Reformulate zero-shot learning as sample-level graph recognition. In: Proceedings of the AAAI conference on artificial intelligence, vol 37, pp 7775\u20137783","DOI":"10.1609\/aaai.v37i6.25942"},{"key":"19915_CR9","doi-asserted-by":"publisher","first-page":"191","DOI":"10.1007\/s11704-017-7031-7","volume":"12","author":"M-L Zhang","year":"2018","unstructured":"Zhang M-L, Li Y-K, Liu X-Y, Geng X (2018) Binary relevance for multi-label learning: an overview. Front Comp Sci 12:191\u2013202","journal-title":"Front Comp Sci"},{"issue":"7","key":"19915_CR10","doi-asserted-by":"publisher","first-page":"2038","DOI":"10.1016\/j.patcog.2006.12.019","volume":"40","author":"M-L Zhang","year":"2007","unstructured":"Zhang M-L, Zhou Z-H (2007) Ml-knn: A lazy learning approach to multi-label learning. Pattern Recogn 40(7):2038\u20132048","journal-title":"Pattern Recogn"},{"issue":"10","key":"19915_CR11","doi-asserted-by":"publisher","first-page":"6642","DOI":"10.1109\/TCSVT.2022.3177320","volume":"32","author":"L Yan","year":"2022","unstructured":"Yan L, Ma S, Wang Q, Chen Y, Zhang X, Savakis A, Liu D (2022) Video captioning using global-local representation. IEEE Trans Circuits Syst Video Technol 32(10):6642\u20136656","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"19915_CR12","doi-asserted-by":"crossref","unstructured":"Liu D, Cui Y, Yan L, Mousas C, Yang B, Chen Y (2021) Densernet: Weakly supervised visual localization using multi-scale feature aggregation. In: Proceedings of the AAAI conference on artificial intelligence, vol 35, pp 6101\u20136109","DOI":"10.1609\/aaai.v35i7.16760"},{"key":"19915_CR13","unstructured":"Radford A, Kim JW, Hallacy C, Ramesh A, Goh G, Agarwal S, Sastry G, Askell A, Mishkin P, Clark J, et al (2021) Learning transferable visual models from natural language supervision. In: International conference on machine learning, PMLR, pp 8748\u20138763"},{"key":"19915_CR14","doi-asserted-by":"crossref","unstructured":"Huynh D, Elhamifar E (2020) A shared multi-attention framework for multi-label zero-shot learning. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 8776\u20138786","DOI":"10.1109\/CVPR42600.2020.00880"},{"key":"19915_CR15","doi-asserted-by":"crossref","unstructured":"Narayan S, Gupta A, Khan S, Khan FS, Shao L, Shah M (2021) Discriminative region-based multi-label zero-shot learning. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 8731\u20138740","DOI":"10.1109\/ICCV48922.2021.00861"},{"key":"19915_CR16","doi-asserted-by":"crossref","unstructured":"Chua T-S, Tang J, Hong R, Li H, Luo Z, Zheng Y (2009) Nus-wide: a real-world web image database from national university of singapore. In: Proceedings of the ACM international conference on image and video retrieval, pp 1\u20139","DOI":"10.1145\/1646396.1646452"},{"issue":"7","key":"19915_CR17","doi-asserted-by":"publisher","first-page":"1956","DOI":"10.1007\/s11263-020-01316-z","volume":"128","author":"A Kuznetsova","year":"2020","unstructured":"Kuznetsova A, Rom H, Alldrin N, Uijlings J, Krasin I, Pont-Tuset J, Kamali S, Popov S, Malloci M, Kolesnikov A et al (2020) The open images dataset v4: Unified image classification, object detection, and visual relationship detection at scale. Int J Comput Vision 128(7):1956\u20131981","journal-title":"Int J Comput Vision"},{"key":"19915_CR18","doi-asserted-by":"crossref","unstructured":"Misra I, Lawrence\u00a0Zitnick C, Mitchell M, Girshick R (2016) Seeing through the human reporting bias: Visual classifiers from noisy human-centric labels. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2930\u20132939","DOI":"10.1109\/CVPR.2016.320"},{"key":"19915_CR19","doi-asserted-by":"crossref","unstructured":"Huynh D, Elhamifar E (2020) Interactive multi-label cnn learning with partial labels. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 9423\u20139432","DOI":"10.1109\/CVPR42600.2020.00944"},{"key":"19915_CR20","doi-asserted-by":"crossref","unstructured":"Wang Y, He D, Li F, Long X, Zhou Z, Ma J, Wen S (2020) Multi-label classification with label graph superimposing. In: Proceedings of the AAAI conference on artificial intelligence, vol 34, pp 12265\u201312272","DOI":"10.1609\/aaai.v34i07.6909"},{"key":"19915_CR21","unstructured":"Weston J, Bengio S, Usunier N (2011) Wsabie: Scaling up to large vocabulary image annotation. In: Twenty-second international joint conference on artificial intelligence"},{"key":"19915_CR22","doi-asserted-by":"crossref","unstructured":"Pennington J, Socher R, Manning CD (2014) Glove: Global vectors for word representation. In: Proceedings of the 2014 conference on empirical methods in natural language processing (EMNLP), pp 1532\u20131543","DOI":"10.3115\/v1\/D14-1162"},{"key":"19915_CR23","doi-asserted-by":"publisher","DOI":"10.1016\/j.patcog.2020.107675","volume":"111","author":"Z Ma","year":"2021","unstructured":"Ma Z, Chen S (2021) Expand globally, shrink locally: Discriminant multi-label learning with missing labels. Pattern Recogn 111:107675","journal-title":"Pattern Recogn"},{"issue":"7","key":"19915_CR24","doi-asserted-by":"publisher","first-page":"1425","DOI":"10.1109\/TPAMI.2015.2487986","volume":"38","author":"Z Akata","year":"2015","unstructured":"Akata Z, Perronnin F, Harchaoui Z, Schmid C (2015) Label-embedding for image classification. IEEE Trans Pattern Anal Mach Intell 38(7):1425\u20131438","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"19915_CR25","unstructured":"Frome A, Corrado GS, Shlens J, Bengio S, Dean J, Ranzato M, Mikolov T (2013) Devise: A deep visual-semantic embedding model. Adv Neural Inform Process Syst 26"},{"key":"19915_CR26","doi-asserted-by":"crossref","unstructured":"Xian Y, Akata Z, Sharma G, Nguyen Q, Hein M, Schiele B (2016) Latent embeddings for zero-shot classification. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 69\u201377","DOI":"10.1109\/CVPR.2016.15"},{"key":"19915_CR27","unstructured":"Mikolov T, Chen K, Corrado G, Dean J (2013) Efficient estimation of word representations in vector space. arXiv preprint arXiv:1301.3781"},{"key":"19915_CR28","doi-asserted-by":"publisher","DOI":"10.1016\/j.knosys.2021.107780","volume":"236","author":"J Liu","year":"2022","unstructured":"Liu J, Fu L, Zhang H, Ye Q, Yang W, Liu L (2022) Learning discriminative and representative feature with cascade gan for generalized zero-shot learning. Knowl-Based Syst 236:107780","journal-title":"Knowl-Based Syst"},{"key":"19915_CR29","unstructured":"Xu B, Zeng Z, Lian C, Ding Z (2022) Generative mixup networks for zero-shot learning. IEEE Trans Neural Netw Learn Syst"},{"key":"19915_CR30","doi-asserted-by":"crossref","unstructured":"Ye Y, Pan T, Luo T, Li J, Shen HT (2022) Learning modality-consistent latent representations for generalized zero-shot learning. IEEE Trans Multimed","DOI":"10.1109\/TMM.2022.3145237"},{"key":"19915_CR31","unstructured":"Ji Z, Fu Y, Guo J, Pang Y, Zhang ZM et al (2018) Stacked semantics-guided attention model for fine-grained zero-shot learning. Adv Neural Inform Process Systems 31"},{"key":"19915_CR32","unstructured":"Gupta A, Narayan S, Khan S, Khan FS, Shao L, Weijer J (2021) Generative multi-label zero-shot learning. arXiv preprint arXiv:2101.11606"},{"key":"19915_CR33","doi-asserted-by":"crossref","unstructured":"Zhang Y, Gong B, Shah M (2016) Fast zero-shot image tagging. In: 2016 IEEE Conference on Computer Vision and Pattern Recognition (CVPR), IEEE, pp 5985\u20135994","DOI":"10.1109\/CVPR.2016.644"},{"key":"19915_CR34","doi-asserted-by":"crossref","unstructured":"Ben-Cohen A, Zamir N, Ben-Baruch E, Friedman I, Zelnik-Manor L (2021) Semantic diversity learning for zero-shot multi-label classification. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 640\u2013650","DOI":"10.1109\/ICCV48922.2021.00068"},{"key":"19915_CR35","unstructured":"Simonyan K, Zisserman A (2014) Very deep convolutional networks for large-scale image recognition. arXiv preprint arXiv:1409.1556"},{"key":"19915_CR36","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 770\u2013778","DOI":"10.1109\/CVPR.2016.90"},{"key":"19915_CR37","doi-asserted-by":"crossref","unstructured":"Zhang Z, Ma W, Wu Y, Wang G (2020) Self-orthogonality module: A network architecture plug-in for learning orthogonal filters. In: Proceedings of the IEEE\/CVF winter conference on applications of computer vision, pp 1050\u20131059","DOI":"10.1109\/WACV45572.2020.9093466"},{"key":"19915_CR38","unstructured":"Norouzi M, Mikolov T, Bengio S, Singer Y, Shlens J, Frome A, Corrado GS, Dean J (2013) Zero-shot learning by convex combination of semantic embeddings. arXiv preprint arXiv:1312.5650"},{"key":"19915_CR39","unstructured":"Kim J-H, Jun J, Zhang B-T (2018) Bilinear attention networks. Advances in neural information processing systems 31"},{"key":"19915_CR40","doi-asserted-by":"crossref","unstructured":"Liang J, Cui Y, Wang Q, Geng T, Wang W, Liu D (2024) Clusterfomer: Clustering as a universal visual learner. Adv Neural Inform Process Syst 36","DOI":"10.1109\/TNNLS.2025.3531987"},{"key":"19915_CR41","doi-asserted-by":"crossref","unstructured":"He S, Guo T, Dai T, Qiao R, Ren B, Xia S-T (2022) Open-vocabulary multi-label classification via multi-modal knowledge transfer. arXiv preprint arXiv:2207.01887","DOI":"10.1609\/aaai.v37i1.25159"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-19915-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-024-19915-0\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-19915-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,7,2]],"date-time":"2025-07-02T11:23:14Z","timestamp":1751455394000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-024-19915-0"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,8,2]]},"references-count":41,"journal-issue":{"issue":"20","published-online":{"date-parts":[[2025,6]]}},"alternative-id":["19915"],"URL":"https:\/\/doi.org\/10.1007\/s11042-024-19915-0","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,8,2]]},"assertion":[{"value":"30 October 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"15 July 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"21 July 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 August 2024","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"We declare that we have no financial and personal relationships with other people or organizations that can inappropriately influence our work, there is no professional or other personal interest of any nature or kind in any product, service and\/or company that could be construed as influencing the position presented in, or the review of, the manuscript entitled, \u201cLabel Correlation Preserving Visual-Semantic Joint Embedding for Multi-Label Zero-Shot Learning\u201d","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest statement"}}]}}