{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,6]],"date-time":"2026-07-06T05:18:17Z","timestamp":1783315097092,"version":"3.54.6"},"reference-count":50,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2026,3,10]],"date-time":"2026-03-10T00:00:00Z","timestamp":1773100800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,3,10]],"date-time":"2026-03-10T00:00:00Z","timestamp":1773100800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s00530-026-02230-z","type":"journal-article","created":{"date-parts":[[2026,3,10]],"date-time":"2026-03-10T13:34:33Z","timestamp":1773149673000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Relative semantic relationship preserving hashing for unsupervised cross-modal retrieval"],"prefix":"10.1007","volume":"32","author":[{"given":"Limeng","family":"Gao","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zhen","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xinzhong","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zhen","family":"Zheng","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Chenchen","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,3,10]]},"reference":[{"key":"2230_CR1","doi-asserted-by":"publisher","unstructured":"Wu, Y., Li, Z.: Mining similarity relationships for unsupervised cross-modal hashing. In: Proceedings of the 2024 IEEE international conference on multimedia and expo, pp. 1\u20136 (2024). https:\/\/doi.org\/10.1109\/ICME57554.2024.10687927","DOI":"10.1109\/ICME57554.2024.10687927"},{"key":"2230_CR2","doi-asserted-by":"publisher","DOI":"10.1016\/j.engappai.2024.108969","author":"X Xie","year":"2024","unstructured":"Xie, X., Li, Z., Li, B., Zhang, C., Ma, H.: Unsupervised cross-modal hashing retrieval via dynamic contrast and optimization. Eng. Appl. Artif. intell. (2024). https:\/\/doi.org\/10.1016\/j.engappai.2024.108969","journal-title":"Eng. Appl. Artif. intell."},{"key":"2230_CR3","doi-asserted-by":"publisher","DOI":"10.1016\/j.patcog.2023.109934","volume":"145","author":"T Yao","year":"2024","unstructured":"Yao, T., Wang, R., Wang, J., Li, Y., Yue, J., Yan, L., Tian, Q.: Efficient supervised graph embedding hashing for large-scale cross-media retrieval. Pattern Recognit. 145, 109934 (2024). https:\/\/doi.org\/10.1016\/j.patcog.2023.109934","journal-title":"Pattern Recognit."},{"issue":"3","key":"2230_CR4","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1007\/s13735-025-00378-4","volume":"14","author":"L Gao","year":"2025","unstructured":"Gao, L., Wang, Z., Wang, X., et al.: Weighted semantic feature based self-supervised deep cross-modal hashing. Int. J. Multimed. Info. Retr. 14(3), 1\u201322 (2025). https:\/\/doi.org\/10.1007\/s13735-025-00378-4","journal-title":"Int. J. Multimed. Info. Retr."},{"key":"2230_CR5","doi-asserted-by":"publisher","DOI":"10.1016\/j.knosys.2024.111837","volume":"295","author":"Z Li","year":"2024","unstructured":"Li, Z., Yao, T., Wang, L., Li, Y., Wang, G.: Supervised contrastive discrete hashing for cross-modal retrieval. Knowl. Based Syst. 295, 111837 (2024). https:\/\/doi.org\/10.1016\/j.knosys.2024.111837","journal-title":"Knowl. Based Syst."},{"issue":"5","key":"2230_CR6","doi-asserted-by":"publisher","first-page":"5091","DOI":"10.1109\/TKDE.2022.3144352","volume":"35","author":"Z Zhang","year":"2023","unstructured":"Zhang, Z., Luo, H., Zhu, L., Lu, G., Shen, H.T.: Modality-invariant asymmetric networks for cross-modal hashing. IEEE Trans. Knowl. Data Eng. 35(5), 5091\u20135104 (2023). https:\/\/doi.org\/10.1109\/TKDE.2022.3144352","journal-title":"IEEE Trans. Knowl. Data Eng."},{"key":"2230_CR7","doi-asserted-by":"publisher","first-page":"9530","DOI":"10.1109\/TMM.2023.3254199","volume":"25","author":"X Liu","year":"2023","unstructured":"Liu, X., Zeng, H., Shi, Y., Zhu, J., Hsia, C., Ma, K.: Deep cross-modal hashing based on semantic consistent ranking. IEEE Trans. Multimedia 25, 9530\u20139542 (2023). https:\/\/doi.org\/10.1109\/TMM.2023.3254199","journal-title":"IEEE Trans. Multimedia"},{"issue":"10","key":"2230_CR8","doi-asserted-by":"publisher","first-page":"7255","DOI":"10.1109\/TCSVT.2022.3172716","volume":"32","author":"Y Shi","year":"2022","unstructured":"Shi, Y., Zhao, Y., Liu, X., Zheng, F., Ou, W., You, X., Peng, Q.: Deep adaptively-enhanced hashing with discriminative similarity guidance for unsupervised cross-modal retrieval. IEEE Trans. Circuits Syst. Video Technol 32(10), 7255\u20137268 (2022). https:\/\/doi.org\/10.1109\/TCSVT.2022.3172716","journal-title":"IEEE Trans. Circuits Syst. Video Technol"},{"key":"2230_CR9","doi-asserted-by":"publisher","unstructured":"Zhang, J., Peng, Y., Yuan, M.: Unsupervised generative adversarial cross-modal hashing. In: Proceedings of the AAAI conference on artificial intelligence, pp. 32 (2018). https:\/\/doi.org\/10.48550\/arXiv.1712.00358","DOI":"10.48550\/arXiv.1712.00358"},{"issue":"9","key":"2230_CR10","doi-asserted-by":"publisher","first-page":"8838","DOI":"10.1109\/TKDE.2022.3218656","volume":"35","author":"L Zhu","year":"2023","unstructured":"Zhu, L., Wu, X., Li, J., Zhang, Z., Guan, W., Shen, H.T.: Work together: correlation-identity reconstruction hashing for unsupervised cross-modal retrieval. IEEE Trans. Knowl. Data Eng. 35(9), 8838\u20138851 (2023). https:\/\/doi.org\/10.1109\/TKDE.2022.3218656","journal-title":"IEEE Trans. Knowl. Data Eng."},{"key":"2230_CR11","doi-asserted-by":"publisher","unstructured":"Wu, G., Lin, Z., Han, J., Liu, L., Ding, G., Zhang, B., Shen, J.: Unsupervised deep hashing via binary latent factor models for large-scale cross-modal retrieval. In: IJCAI, pp. 2854\u20132860 (2018). https:\/\/doi.org\/10.3724\/SP.J.1089.2021.18599","DOI":"10.3724\/SP.J.1089.2021.18599"},{"key":"2230_CR12","doi-asserted-by":"crossref","unstructured":"Yu, J., Zhou, H., Zhan, Y., Tao, D.: Deep graph-neighbor coherence preserving network for unsupervised cross-modal hashing. In: Proceedings of the AAAI conference on artificial intelligence, vol. 35, pp. 4626\u20134634 (2021)","DOI":"10.1609\/aaai.v35i5.16592"},{"key":"2230_CR13","doi-asserted-by":"publisher","unstructured":"Yang, D., Wu, D., Zhang, W., Zhang, H., Li, B., Wang, W.: Deep semantic-alignment hashing for unsupervised cross-modal retrieval. In: Proceedings of the 2020 international conference on multimedia retrieval, pp. 44\u201352 (2020). https:\/\/doi.org\/10.1145\/3372278.3390673","DOI":"10.1145\/3372278.3390673"},{"issue":"10","key":"2230_CR14","doi-asserted-by":"publisher","first-page":"7255","DOI":"10.1109\/TCSVT.2022.3172716","volume":"32","author":"Y Shi","year":"2022","unstructured":"Shi, Y., Zhao, Y., Liu, X., Zheng, F., Ou, W., You, X., Peng, Q.: Deep adaptively-enhanced hashing with discriminative similarity guidance for unsupervised cross-modal retrieval. IEEE Trans. Circuits Syst. Video Technol. 32(10), 7255\u201368 (2022). https:\/\/doi.org\/10.1109\/TCSVT.2022.3172716","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"2230_CR15","doi-asserted-by":"publisher","DOI":"10.1007\/s13042-024-02154-y","author":"S Xiong","year":"2024","unstructured":"Xiong, S., Pan, L., Ma, X., Hu, Q., Beckman, E.: Unsupervised deep hashing with multiple similarity preservation for cross-modal image-text retrieval. Int. J. Mach. Learn. Cybern. (2024). https:\/\/doi.org\/10.1007\/s13042-024-02154-y","journal-title":"Int. J. Mach. Learn. Cybern."},{"key":"2230_CR16","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.127911","volume":"595","author":"H Meng","year":"2024","unstructured":"Meng, H., Zhang, H., Liu, L., Liu, D., Lu, X., Guo, X.: Joint-modal graph convolutional hashing for unsupervised cross-modal retrieval. Neurocomputing , 127911 (2024). https:\/\/doi.org\/10.1016\/j.neucom.2024.127911","journal-title":"Neurocomputing"},{"key":"2230_CR17","doi-asserted-by":"publisher","unstructured":"Qiu, Z., Su, Q., Ou, Z., Yu, J., Chen, C.: Unsupervised hashing with contrastive information bottleneck. In: IJCAI, pp. 959\u2013965 (2021). https:\/\/doi.org\/10.1109\/ICASSP43922.2022.9746251","DOI":"10.1109\/ICASSP43922.2022.9746251"},{"key":"2230_CR18","doi-asserted-by":"publisher","unstructured":"Mikriukov, G., Ravanbakhsh, M., Demir, B.: Deep unsupervised contrastive hashing for large-scale cross-modal retrieval in remote sensing. In: 2022 IEEE international conference on acoustics, speech, and signal processing, pp. 4463\u20134467 (2022). https:\/\/doi.org\/10.48550\/arXiv.2201.08125","DOI":"10.48550\/arXiv.2201.08125"},{"issue":"3","key":"2230_CR19","doi-asserted-by":"publisher","first-page":"3877","DOI":"10.1109\/TPAMI.2022.3177356","volume":"45","author":"P Hu","year":"2022","unstructured":"Hu, P., Zhu, H., Lin, J., Peng, D., Zhao, Y.P., Peng, X.: Unsupervised contrastive cross-modal hashing. IEEE Trans. Pattern Anal. Mach. Intell. 45(3), 3877\u20133889 (2022). https:\/\/doi.org\/10.1109\/TPAMI.2022.3177356","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"2230_CR20","doi-asserted-by":"publisher","unstructured":"Lu, K., Yu, Y., Liang, M., Zhang, M., Cao, X., Zhao, Z., Yin, M., Xue, Z.: Deep unsupervised momentum contrastive hashing for cross-modal retrieval. In: 2023 IEEE international conference on multimedia and expo (ICME), pp. 126\u2013131 (2023). https:\/\/doi.org\/10.1109\/ICME55011.2023.00030","DOI":"10.1109\/ICME55011.2023.00030"},{"key":"2230_CR21","doi-asserted-by":"publisher","DOI":"10.1016\/j.ins.2023.119543","author":"C Xie","year":"2023","unstructured":"Xie, C., Gao, Y., Zhou, Q., Zhou, J.: Multi-similarity reconstructing and clustering-based contrastive hashing for cross-modal retrieval. Inf. Sci. (2023). https:\/\/doi.org\/10.1016\/j.ins.2023.119543","journal-title":"Inf. Sci."},{"issue":"7","key":"2230_CR22","doi-asserted-by":"publisher","first-page":"3490","DOI":"10.1109\/TIP.2019.2897944","volume":"28","author":"Q Jiang","year":"2019","unstructured":"Jiang, Q., Li, W.: Discrete latent factor model for cross-modal hashing. IEEE Trans. Image Process. 28(7), 3490\u20133501 (2019). https:\/\/doi.org\/10.1109\/TIP.2019.2897944","journal-title":"IEEE Trans. Image Process."},{"issue":"11","key":"2230_CR23","doi-asserted-by":"publisher","first-page":"3507","DOI":"10.1109\/tkde.2020.2974825","volume":"33","author":"Y Wang","year":"2020","unstructured":"Wang, Y., Luo, X., Nie, L., Song, J., Zhang, W., Xu, X.: BATCH: a scalable asymmetric discrete cross-modal hashing. IEEE Trans. Knowl. Data Eng. 33(11), 3507\u20133519 (2020). https:\/\/doi.org\/10.1109\/tkde.2020.2974825","journal-title":"IEEE Trans. Knowl. Data Eng."},{"issue":"2","key":"2230_CR24","doi-asserted-by":"publisher","first-page":"1365","DOI":"10.1109\/TKDE.2021.3099125","volume":"35","author":"D Zhang","year":"2021","unstructured":"Zhang, D., Wu, X.J., Xu, T., Yin, H.: DAH: discrete asymmetric hashing for efficient cross-media retrieval. IEEE Trans. Knowl. Data Eng. 35(2), 1365\u20131378 (2021). https:\/\/doi.org\/10.1109\/TKDE.2021.3099125","journal-title":"IEEE Trans. Knowl. Data Eng."},{"issue":"6","key":"2230_CR25","doi-asserted-by":"publisher","first-page":"6461","DOI":"10.1109\/TKDE.2022.3159131","volume":"35","author":"D Zhang","year":"2022","unstructured":"Zhang, D., Wu, X.J., Xu, T., Kittler, J.: Watch: two-stage discrete cross-media hashing. IEEE Trans. Knowl. Data Eng. 35(6), 6461\u20136474 (2022). https:\/\/doi.org\/10.1109\/TKDE.2022.3159131","journal-title":"IEEE Trans. Knowl. Data Eng."},{"key":"2230_CR26","doi-asserted-by":"publisher","unstructured":"Jiang, Q., Li, W.: Deep cross-modal hashing. In: Proceedings of the computer vision and pattern recognition (CVPR), pp. 3270\u20133278 (2017). https:\/\/doi.org\/10.1109\/CVPR.2017.348","DOI":"10.1109\/CVPR.2017.348"},{"issue":"4","key":"2230_CR27","doi-asserted-by":"publisher","first-page":"1838","DOI":"10.1109\/TNNLS.2020.2997020","volume":"34","author":"L Jin","year":"2023","unstructured":"Jin, L., Li, Z., Tang, J.: Deep semantic multimodal hashing network for scalable image-text and video-text retrievals. IEEE Trans. Neural Netw. Learn. Syst. 34(4), 1838\u20131851 (2023). https:\/\/doi.org\/10.1109\/TNNLS.2020.2997020","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"key":"2230_CR28","doi-asserted-by":"publisher","DOI":"10.1016\/j.ins.2023.120064","author":"C Zheng","year":"2024","unstructured":"Zheng, C., Zhu, L., Zhang, Z., Duan, W., Lu, W.: LCEMH: label correlation enhanced multi-modal hashing for efficient multi-modal retrieval. Inf. Sci. (2024). https:\/\/doi.org\/10.1016\/j.ins.2023.120064","journal-title":"Inf. Sci."},{"issue":"6","key":"2230_CR29","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3643639","volume":"20","author":"Y Huo","year":"2024","unstructured":"Huo, Y., Qin, Q., Dai, J., Zhang, W., Huang, L., Wang, C.: Deep neighborhood-aware proxy hashing with uniform distribution constraint for cross-modal retrieval. ACM Trans. Multimed. Comput. Commun. Appl. 20(6), 1\u201323 (2024). https:\/\/doi.org\/10.1145\/3643639","journal-title":"ACM Trans. Multimed. Comput. Commun. Appl."},{"key":"2230_CR30","doi-asserted-by":"publisher","unstructured":"Kumar, S., Udupa, R.: Learning hash functions for cross-view similarity search. In: Twenty-second international joint conference on artifcial intelligence (2011). https:\/\/doi.org\/10.5591\/978-1-57735-516-8\/IJCAI11-230","DOI":"10.5591\/978-1-57735-516-8\/IJCAI11-230"},{"key":"2230_CR31","doi-asserted-by":"publisher","unstructured":"Song, J., Yang, Y., Yang, Y., Huang, Z., Shen, H.T.: Inter-media hashing for large-scale retrieval from heterogeneous data sources. In: Proceedings of the ACM SIGMOD international conference on management of data, pp. 785\u2013796 (2013). https:\/\/doi.org\/10.1145\/2463676.2465274","DOI":"10.1145\/2463676.2465274"},{"key":"2230_CR32","doi-asserted-by":"publisher","unstructured":"Zhou, J., Ding, G., Guo, Y.: Latent semantic sparse hashing for cross-modal similarity search. In: Proceedings of the 37th international ACM SIGIR conference on research & development in information retrieval, pp. 415\u2013424 (2014). https:\/\/doi.org\/10.1145\/2600428.2609610","DOI":"10.1145\/2600428.2609610"},{"key":"2230_CR33","doi-asserted-by":"publisher","unstructured":"Ding, G., Guo, Y., Zhou, J.: Collective matrix factorization hashing for multimodal data. In: Proceedings of the 2014 IEEE conference on computer vision and pattern recognition (CVPR), pp. 2083\u20132090 (2014). https:\/\/doi.org\/10.1109\/CVPR.2014.267","DOI":"10.1109\/CVPR.2014.267"},{"issue":"4","key":"2230_CR34","doi-asserted-by":"publisher","first-page":"973","DOI":"10.1145\/3123266.3123355","volume":"21","author":"X Li","year":"2018","unstructured":"Li, X., Hu, D., Nie, F.: Deep binary reconstruction for cross-modal hashing. IEEE Trans. Multimedia 21(4), 973\u2013985 (2018). https:\/\/doi.org\/10.1145\/3123266.3123355","journal-title":"IEEE Trans. Multimedia"},{"key":"2230_CR35","doi-asserted-by":"publisher","unstructured":"Su, S., Zhong, Z., Zhang, C.: Deep joint-semantics reconstructing hashing for large-scale unsupervised cross-modal retrieval. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp. 3027\u20133035 (2019). https:\/\/doi.org\/10.1109\/ICCV.2019.00312","DOI":"10.1109\/ICCV.2019.00312"},{"key":"2230_CR36","doi-asserted-by":"publisher","unstructured":"Liu, S., Qian, S., Guan, Y., Zhan, J., Ying, L.: Joint-modal distribution-based similarity hashing for large-scale unsupervised deep cross-modal retrieval. In: Proceedings of the 43rd international ACM SIGIR conference on research and development in information retrieval, pp. 1379\u20131388 (2020). https:\/\/doi.org\/10.1145\/3397271.3401086","DOI":"10.1145\/3397271.3401086"},{"key":"2230_CR37","doi-asserted-by":"publisher","first-page":"466","DOI":"10.1109\/TMM.2021.3053766","volume":"24","author":"PF Zhang","year":"2021","unstructured":"Zhang, P.F., Li, Y., Huang, Z., Xu, X.S.: Aggregation-based graph convolutional hashing for unsupervised cross-modal retrieval. IEEE Trans. Multimedia 24, 466\u2013479 (2021). https:\/\/doi.org\/10.1109\/TMM.2021.3053766","journal-title":"IEEE Trans. Multimedia"},{"key":"2230_CR38","doi-asserted-by":"publisher","DOI":"10.1016\/j.inffus.2023.101968","volume":"100","author":"X Xia","year":"2023","unstructured":"Xia, X., Dong, G., Li, F., Zhu, L., Ying, X.: When clip meets cross-modal hashing retrieval: a new strong baseline. Inf. Fusion 100, 101968 (2023). https:\/\/doi.org\/10.1016\/j.inffus.2023.101968","journal-title":"Inf. Fusion"},{"key":"2230_CR39","doi-asserted-by":"publisher","DOI":"10.1016\/j.neunet.2024.106211","volume":"174","author":"J Cui","year":"2024","unstructured":"Cui, J., He, Z., Huang, Q., Fu, Y., Li, Y., Wen, J.: Structure-aware contrastive hashing for unsupervised cross-modal retrieval. Neural Netw 174, 106211 (2024). https:\/\/doi.org\/10.1016\/j.neunet.2024.106211","journal-title":"Neural Netw"},{"key":"2230_CR40","doi-asserted-by":"publisher","DOI":"10.1016\/j.eswa.2023.121516","volume":"237","author":"D Yao","year":"2024","unstructured":"Yao, D., Li, Z., Li, B., Zhang, C., Ma, H.: Similarity graph-correlation reconstruction network for unsupervised cross-modal hashing. Expert Syst. Appl. 237, 121516 (2024). https:\/\/doi.org\/10.1016\/j.eswa.2023.121516","journal-title":"Expert Syst. Appl."},{"issue":"1","key":"2230_CR41","doi-asserted-by":"publisher","DOI":"10.1016\/j.ipm.2024.103958","volume":"62","author":"Y Chen","year":"2025","unstructured":"Chen, Y., Long, Y., Yang, Z., Long, J.: Unsupervised adaptive hypergraph correlation hashing for multimedia retrieval. Inf. Process. Manag. 62(1), 103958 (2025). https:\/\/doi.org\/10.1016\/j.ipm.2024.103958","journal-title":"Inf. Process. Manag."},{"key":"2230_CR42","doi-asserted-by":"publisher","DOI":"10.1016\/j.neunet.2024.106923","volume":"182","author":"Y Chen","year":"2025","unstructured":"Chen, Y., Long, Y., Yang, Z., Long, J.: Parameter adaptive contrastive hashing for multimedia retrieval. Neural Netw 182, 106923 (2025). https:\/\/doi.org\/10.1016\/j.neunet.2024.106923","journal-title":"Neural Netw"},{"key":"2230_CR43","doi-asserted-by":"publisher","first-page":"2","DOI":"10.1007\/s13735-023-00268-7","volume":"12","author":"M Li","year":"2023","unstructured":"Li, M., Li, Y., Ge, M., Ma, L.: CLIP-based fusion-modal reconstructing hashing for large-scale unsupervised cross-modal retrieval. Int J Multimed Info Retr. 12, 2 (2023). https:\/\/doi.org\/10.1007\/s13735-023-00268-7","journal-title":"Int J Multimed Info Retr."},{"key":"2230_CR44","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.128844","volume":"614","author":"Y Wu","year":"2025","unstructured":"Wu, Y., Li, B., Li, Z.: Revising similarity relationship hashing for unsupervised cross-modal retrieval. Neurocomputing 614, 128844 (2025). https:\/\/doi.org\/10.1016\/j.neucom.2024.128844","journal-title":"Neurocomputing"},{"key":"2230_CR45","doi-asserted-by":"publisher","DOI":"10.1007\/s00530-024-01539-x","volume":"30","author":"H Kang","year":"2024","unstructured":"Kang, H., Zhang, X., Han, W., Zhou, M.: Dark knowledge association guided hashing for unsupervised cross-modal retrieval. Multimedia Syst. 30, 375 (2024). https:\/\/doi.org\/10.1007\/s00530-024-01539-x","journal-title":"Multimedia Syst."},{"key":"2230_CR46","unstructured":"Kingma, D.P., Ba, J.: Adam: a method for stochastic optimization. arXiv preprint arXiv:1412.6980 (2014)"},{"key":"2230_CR47","doi-asserted-by":"publisher","unstructured":"Huiskes, M.J., Lew, M.S.: The mir fickr retrieval evaluation. In: Proceedings of the 1st ACM international conference on multimedia information retrieval, pp. 39\u201343 (2008). https:\/\/doi.org\/10.1145\/1460096.1460104","DOI":"10.1145\/1460096.1460104"},{"key":"2230_CR48","doi-asserted-by":"publisher","unstructured":"Chua, T.-S., Tang, J., Hong, R., Li, H., Luo, Z., Zheng, Y.: Nus-wide: a real-world web image database from National University of Singapore. In: Proceedings of the ACM international conference on image and video retrieval, pp. 1\u20139 (2009). https:\/\/doi.org\/10.1145\/1646396.1646452","DOI":"10.1145\/1646396.1646452"},{"key":"2230_CR49","doi-asserted-by":"publisher","unstructured":"Lin, T.-Y., Maire, M., Belongie, S., Hays, J., Perona, P., Ramanan, D., Doll\u00e1r, P., Zitnick, C.L.: Microsoft coco: common objects in context. In: Computer vision\u2013ECCV 2014: 13th European conference, Zurich, Switzerland, September 6\u201312, 2014, proceedings, part V 13, pp. 740\u2013755 (2014). https:\/\/doi.org\/10.1007\/978-3-319-10602-1_48","DOI":"10.1007\/978-3-319-10602-1_48"},{"key":"2230_CR50","doi-asserted-by":"publisher","unstructured":"Radford, A., Kim, J.W., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell. A., Mishkin, P., Clark, J., et\u00a0al.: Learning transferable visual models from natural language supervision. In: International conference on machine learning, pp. 8748\u20138763. PMLR (2021). https:\/\/doi.org\/10.48550\/arXiv.2103.00020","DOI":"10.48550\/arXiv.2103.00020"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02230-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-026-02230-z","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02230-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,6]],"date-time":"2026-07-06T05:07:46Z","timestamp":1783314466000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-026-02230-z"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,3,10]]},"references-count":50,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2230"],"URL":"https:\/\/doi.org\/10.1007\/s00530-026-02230-z","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,3,10]]},"assertion":[{"value":"15 May 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"13 January 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"10 March 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"185"}}