{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,9,26]],"date-time":"2025-09-26T11:10:14Z","timestamp":1758885014120,"version":"3.44.0"},"reference-count":39,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2025,6,26]],"date-time":"2025-06-26T00:00:00Z","timestamp":1750896000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,6,26]],"date-time":"2025-06-26T00:00:00Z","timestamp":1750896000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"National Natural Science Foundation of China under Grants","award":["62376017"],"award-info":[{"award-number":["62376017"]}]},{"DOI":"10.13039\/501100012226","name":"Fundamental Research Funds for the Central Universities","doi-asserted-by":"crossref","award":["buctrc202221"],"award-info":[{"award-number":["buctrc202221"]}],"id":[{"id":"10.13039\/501100012226","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Multimed Info Retr"],"published-print":{"date-parts":[[2025,9]]},"DOI":"10.1007\/s13735-025-00371-x","type":"journal-article","created":{"date-parts":[[2025,6,26]],"date-time":"2025-06-26T03:51:50Z","timestamp":1750909910000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["MCDINO: Self-supervised learning of masks based on combination of multi-path channel attention and local feature weighting"],"prefix":"10.1007","volume":"14","author":[{"given":"Yunxue","family":"Shao","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zhiyang","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Lingfeng","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,6,26]]},"reference":[{"issue":"1","key":"371_CR1","doi-asserted-by":"publisher","first-page":"255","DOI":"10.1146\/annurev.cs.04.060190.001351","volume":"4","author":"TG Dietterich","year":"1990","unstructured":"Dietterich TG (1990) Machine learning. Annual review of computer science 4(1):255\u2013306","journal-title":"Annual review of computer science"},{"key":"371_CR2","unstructured":"Mitchell TM, Mitchell TM (1997) Machine learning, volume\u00a01. McGraw-hill New York,"},{"key":"371_CR3","doi-asserted-by":"crossref","unstructured":"Zhou Z-H (2021) Machine learning. Springer nature","DOI":"10.1007\/978-981-15-1967-3"},{"issue":"1","key":"371_CR4","first-page":"857","volume":"35","author":"X Liu","year":"2021","unstructured":"Liu X, Zhang F, Hou Z, Mian L, Wang Z, Zhang J, Tang J (2021) Self-supervised learning: Generative or contrastive. IEEE Trans Knowl Data Eng 35(1):857\u2013876","journal-title":"IEEE Trans Knowl Data Eng"},{"issue":"1","key":"371_CR5","doi-asserted-by":"publisher","first-page":"2","DOI":"10.3390\/technologies9010002","volume":"9","author":"A Jaiswal","year":"2020","unstructured":"Jaiswal A, Babu AR, Zadeh MZ, Banerjee D, Makedon F (2020) A survey on contrastive self-supervised learning. Technologies 9(1):2","journal-title":"Technologies"},{"key":"371_CR6","unstructured":"Ballard DH, Brown CM (1982) Computer vision. Prentice Hall Professional Technical Reference"},{"issue":"9","key":"371_CR7","doi-asserted-by":"publisher","first-page":"6186","DOI":"10.1109\/TCSVT.2022.3162599","volume":"32","author":"J Nie","year":"2022","unstructured":"Nie J, Han W, He Z, Gao M, Dong Z (2022) Spreading fine-grained prior knowledge for accurate tracking. IEEE Trans Circuits Syst Video Technol 32(9):6186\u20136199","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"371_CR8","doi-asserted-by":"publisher","first-page":"6194","DOI":"10.1109\/TMM.2022.3206668","volume":"25","author":"J Nie","year":"2022","unstructured":"Nie J, He Z, Yang Y, Gao M, Dong Z (2022) Learning localization-aware target confidence for siamese visual tracking. IEEE Trans Multimedia 25:6194\u20136206","journal-title":"IEEE Trans Multimedia"},{"key":"371_CR9","doi-asserted-by":"publisher","first-page":"529","DOI":"10.1016\/j.ins.2023.02.083","volume":"632","author":"J Nie","year":"2023","unstructured":"Nie J, Dong Z, He Z, Han W, Gao M (2023) Faml-rt: Feature alignment-based multi-level similarity metric learning network for a two-stage robust tracker. Inf Sci 632:529\u2013542","journal-title":"Inf Sci"},{"key":"371_CR10","doi-asserted-by":"crossref","unstructured":"Liu B, Liu B (2011) Supervised learning. Web Data Mining: Exploring Hyperlinks, Contents, and Usage Data, pages 63\u2013132","DOI":"10.1007\/978-3-642-19460-3_3"},{"issue":"4","key":"371_CR11","doi-asserted-by":"publisher","first-page":"2761","DOI":"10.1007\/s11831-023-09884-2","volume":"30","author":"V Rani","year":"2023","unstructured":"Rani V, Nabi ST, Kumar M, Mittal A, Kumar K (2023) Self-supervised learning: A succinct review. Archives of Computational Methods in Engineering 30(4):2761\u20132775","journal-title":"Archives of Computational Methods in Engineering"},{"key":"371_CR12","doi-asserted-by":"crossref","unstructured":"Gui J, Chen T, Zhang J, Cao Q, Sun Z, Luo H, Tao D (2024) A survey on self-supervised learning: Algorithms, applications, and future trends. IEEE Trans Pattern Anal Mach Intell","DOI":"10.1109\/TPAMI.2024.3415112"},{"key":"371_CR13","doi-asserted-by":"crossref","unstructured":"Caron M, Touvron H, Misra I, J\u00e9gou H, Mairal J, Bojanowski P, Joulin A (2021) Emerging properties in self-supervised vision transformers. In Proceedings of the IEEE\/CVF international conference on computer vision pages 9650\u20139660","DOI":"10.1109\/ICCV48922.2021.00951"},{"issue":"1","key":"371_CR14","doi-asserted-by":"publisher","first-page":"87","DOI":"10.1109\/TPAMI.2022.3152247","volume":"45","author":"K Han","year":"2022","unstructured":"Han K, Wang Y, Chen H, Chen X, Guo J, Liu Z, Tang Y, Xiao A, Chunjing X, Yixing X et al (2022) A survey on vision transformer. IEEE Trans Pattern Anal Mach Intell 45(1):87\u2013110","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"issue":"3","key":"371_CR15","doi-asserted-by":"publisher","first-page":"331","DOI":"10.1007\/s41095-022-0271-y","volume":"8","author":"M-H Guo","year":"2022","unstructured":"Guo M-H, Tian-Xing X, Liu J-J, Liu Z-N, Jiang P-T, Tai-Jiang M, Zhang S-H, Martin RR, Cheng M-M, Shi-Min H (2022) Attention mechanisms in computer vision: A survey. Computational visual media 8(3):331\u2013368","journal-title":"Computational visual media"},{"key":"371_CR16","doi-asserted-by":"publisher","first-page":"48","DOI":"10.1016\/j.neucom.2021.03.091","volume":"452","author":"Z Niu","year":"2021","unstructured":"Niu Z, Zhong G, Hui Yu (2021) A review on the attention mechanism of deep learning. Neurocomputing 452:48\u201362","journal-title":"Neurocomputing"},{"key":"371_CR17","doi-asserted-by":"crossref","unstructured":"Deng J, Dong W, Socher R, Li L-J, Li K, Fei-Fei L (2009) Imagenet: A large-scale hierarchical image database. In 2009 IEEE conference on computer vision and pattern recognition, pages 248\u2013255. Ieee","DOI":"10.1109\/CVPR.2009.5206848"},{"key":"371_CR18","unstructured":"Beyer L, H\u00e9naff OJ, Kolesnikov A, Zhai X, van\u00a0den Oord A (2020) Are we done with imagenet? arXiv preprint arXiv:2006.07159"},{"issue":"7553","key":"371_CR19","first-page":"436","volume":"521","author":"Y LeCun","year":"2015","unstructured":"LeCun Y, Bengio Y, Hinton G (2015) Deep learning. nature 521(7553):436\u2013444","journal-title":"Deep learning. nature"},{"key":"371_CR20","unstructured":"Ian Goodfellow (2016) Deep learning"},{"issue":"5","key":"371_CR21","doi-asserted-by":"publisher","first-page":"544","DOI":"10.1136\/amiajnl-2011-000464","volume":"18","author":"PM Nadkarni","year":"2011","unstructured":"Nadkarni PM, Ohno-Machado L, Chapman WW (2011) Natural language processing: an introduction. J Am Med Inform Assoc 18(5):544\u2013551","journal-title":"J Am Med Inform Assoc"},{"key":"371_CR22","doi-asserted-by":"crossref","unstructured":"KR1442 Chowdhary and KR Chowdhary. (2020) Natural language processing. Fundamentals of artificial intelligence pages 603\u2013649","DOI":"10.1007\/978-81-322-3972-7_19"},{"key":"371_CR23","doi-asserted-by":"crossref","unstructured":"Fanni SC, Febi M, Aghakhanyan G, Neri E (2023) Natural language processing. In Introduction to Artificial Intelligence, pages 87\u201399. Springer","DOI":"10.1007\/978-3-031-25928-9_5"},{"key":"371_CR24","unstructured":"Alaparthi S, Mishra M (2020) Bidirectional encoder representations from transformers (bert): A sentiment analysis odyssey. arXiv preprint arXiv:2007.01127"},{"issue":"2","key":"371_CR25","doi-asserted-by":"publisher","first-page":"373","DOI":"10.1007\/s10994-019-05855-6","volume":"109","author":"JE Van Engelen","year":"2020","unstructured":"Van Engelen JE, Hoos HH (2020) A survey on semi-supervised learning. Mach Learn 109(2):373\u2013440","journal-title":"Mach Learn"},{"key":"371_CR26","doi-asserted-by":"crossref","unstructured":"Zhai X, Oliver A, Kolesnikov A, Beyer L (2019) S4l: Self-supervised semi-supervised learning. In Proceedings of the IEEE\/CVF international conference on computer vision pages 1476\u20131485","DOI":"10.1109\/ICCV.2019.00156"},{"key":"371_CR27","first-page":"22243","volume":"33","author":"T Chen","year":"2020","unstructured":"Chen T, Kornblith S, Swersky K, Norouzi M, Hinton GE (2020) Big self-supervised models are strong semi-supervised learners. Adv Neural Inf Process Syst 33:22243\u201322255","journal-title":"Adv Neural Inf Process Syst"},{"key":"371_CR28","unstructured":"Hall E (1979) Computer image processing and recognition. Elsevier"},{"key":"371_CR29","doi-asserted-by":"crossref","unstructured":"Acharya T (2005) Image Processing-Principles and Applications. Wiley-Interscience","DOI":"10.1002\/0471745790"},{"issue":"10","key":"371_CR30","doi-asserted-by":"publisher","first-page":"2279","DOI":"10.1016\/S0031-3203(01)00178-9","volume":"35","author":"M Egmont-Petersen","year":"2002","unstructured":"Egmont-Petersen M, de Ridder D, Handels H (2002) Image processing with neural networks-a review. Pattern Recogn 35(10):2279\u20132301","journal-title":"Pattern Recogn"},{"key":"371_CR31","unstructured":"Liu Y (2019) Roberta: A robustly optimized bert pretraining approach. arXiv preprint arXiv:1907.11692, 364"},{"key":"371_CR32","doi-asserted-by":"crossref","unstructured":"He K, Chen X, Xie S, Li Y, Doll\u00e1r P, Girshick R (2022) Masked autoencoders are scalable vision learners. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition pages 16000\u201316009","DOI":"10.1109\/CVPR52688.2022.01553"},{"key":"371_CR33","unstructured":"Bao H, Dong L, Piao S, Wei F (2021) Beit: Bert pre-training of image transformers. arXiv preprint arXiv:2106.08254"},{"key":"371_CR34","doi-asserted-by":"crossref","unstructured":"Xie Z, Zhang Z, Cao Y, Lin Y, Bao J, Yao Z, Dai Q, Han H (2022) Simmim: A simple framework for masked image modeling. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition pages 9653\u20139663","DOI":"10.1109\/CVPR52688.2022.00943"},{"key":"371_CR35","doi-asserted-by":"crossref","unstructured":"Wei C, Fan H, Xie S, Chao-Yuan W, Yuille A, Feichtenhofer C (2022) Masked feature prediction for self-supervised visual pre-training. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition pages 4668\u201314678","DOI":"10.1109\/CVPR52688.2022.01426"},{"key":"371_CR36","doi-asserted-by":"crossref","unstructured":"Xia Z, Pan X, Song S, Li LE, Huang G (2022) Vision transformer with deformable attention. In Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition pages 4794\u20134803","DOI":"10.1109\/CVPR52688.2022.00475"},{"key":"371_CR37","doi-asserted-by":"crossref","unstructured":"Radenovi\u0107 F, Iscen A, Tolias G, Avrithis Y, Chum O (2018) Revisiting oxford and paris: Large-scale image retrieval benchmarking. In Proceedings of the IEEE conference on computer vision and pattern recognition pages 5706\u20135715","DOI":"10.1109\/CVPR.2018.00598"},{"key":"371_CR38","doi-asserted-by":"crossref","unstructured":"Philbin J, Chum O, Isard M, Sivic J, Zisserman A (2008) Lost in quantization: Improving particular object retrieval in large scale image databases. In 2008 IEEE conference on computer vision and pattern recognition, pages 1\u20138. IEEE","DOI":"10.1109\/CVPR.2008.4587635"},{"key":"371_CR39","doi-asserted-by":"crossref","unstructured":"Douze M, J\u00e9gou H, Sandhawalia H, Amsaleg L, Schmid C (2009) Evaluation of gist descriptors for web-scale image search. In Proceedings of the ACM international conference on image and video retrieval pages 1\u20138","DOI":"10.1145\/1646396.1646421"}],"container-title":["International Journal of Multimedia Information Retrieval"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13735-025-00371-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s13735-025-00371-x\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13735-025-00371-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,26]],"date-time":"2025-09-26T10:49:58Z","timestamp":1758883798000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s13735-025-00371-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,6,26]]},"references-count":39,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,9]]}},"alternative-id":["371"],"URL":"https:\/\/doi.org\/10.1007\/s13735-025-00371-x","relation":{},"ISSN":["2192-6611","2192-662X"],"issn-type":[{"type":"print","value":"2192-6611"},{"type":"electronic","value":"2192-662X"}],"subject":[],"published":{"date-parts":[[2025,6,26]]},"assertion":[{"value":"5 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 April 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 June 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"26 June 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of Interest"}}],"article-number":"25"}}