{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,6,6]],"date-time":"2025-06-06T09:10:57Z","timestamp":1749201057072,"version":"3.37.3"},"reference-count":42,"publisher":"Springer Science and Business Media LLC","issue":"7","license":[{"start":{"date-parts":[[2021,11,12]],"date-time":"2021-11-12T00:00:00Z","timestamp":1636675200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"},{"start":{"date-parts":[[2021,11,12]],"date-time":"2021-11-12T00:00:00Z","timestamp":1636675200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["61872191","41571389"],"award-info":[{"award-number":["61872191","41571389"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Neural Comput &amp; Applic"],"published-print":{"date-parts":[[2022,4]]},"DOI":"10.1007\/s00521-021-06696-y","type":"journal-article","created":{"date-parts":[[2021,11,12]],"date-time":"2021-11-12T09:04:13Z","timestamp":1636707853000},"page":"5397-5416","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":8,"title":["Deep semantic hashing with dual attention for cross-modal retrieval"],"prefix":"10.1007","volume":"34","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-3109-2553","authenticated-orcid":false,"given":"Jiagao","family":"Wu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Weiwei","family":"Weng","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Junxia","family":"Fu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Linfeng","family":"Liu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Bin","family":"Hu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2021,11,12]]},"reference":[{"issue":"8","key":"6696_CR1","doi-asserted-by":"publisher","first-page":"3893","DOI":"10.1109\/TIP.2018.2821921","volume":"27","author":"C Deng","year":"2018","unstructured":"Deng C, Chen Z, Liu X, Gao X, Tao D (2018) Triplet-based deep hashing network for cross-modal retrieval. IEEE Trans Image Process 27(8):3893","journal-title":"IEEE Trans Image Process"},{"key":"6696_CR2","doi-asserted-by":"crossref","unstructured":"Cao Y, Long M, Wang J, Liu S (2017) Collective deep quantization for efficient cross-modal retrieval. In: 31st AAAI conference on artificial intelligence, pp 3974\u20133980","DOI":"10.1609\/aaai.v31i1.11218"},{"key":"6696_CR3","doi-asserted-by":"crossref","unstructured":"Wang B, Yang Y, Xu X, Hanjalic A, Shen H (2017) Adversarial cross-modal retrieval. In: Proceedings of the 2017 ACM multimedia conference, pp 154\u2013162","DOI":"10.1145\/3123266.3123326"},{"key":"6696_CR4","doi-asserted-by":"crossref","unstructured":"Wu Y, Wang S, Huang Q (2017) Online asymmetric similarity learning for cross-modal retrieval. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 4269\u20134278","DOI":"10.1109\/CVPR.2017.424"},{"issue":"3","key":"6696_CR5","doi-asserted-by":"publisher","first-page":"370","DOI":"10.1109\/TMM.2015.2390499","volume":"17","author":"C Kang","year":"2015","unstructured":"Kang C, Xiang S, Liao S, Xu C, Pan C (2015) Learning consistent feature representation for cross-modal multimedia retrieval. IEEE Trans Multimedia 17(3):370","journal-title":"IEEE Trans Multimedia"},{"key":"6696_CR6","doi-asserted-by":"crossref","unstructured":"Zhang Y, Lu H (2018) Deep cross-modal projection learning for image-text matching. In: Proceedings of the European Conference on Computer Vision, pp 686\u2013701","DOI":"10.1007\/978-3-030-01246-5_42"},{"key":"6696_CR7","doi-asserted-by":"crossref","unstructured":"Gu J, Cai J, Joty S.R., Niu L, Wang G (2018) Look, imagine and match: Improving textual-visual cross-modal retrieval with generative models. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 7181\u20137189","DOI":"10.1109\/CVPR.2018.00750"},{"issue":"4","key":"6696_CR8","doi-asserted-by":"publisher","first-page":"973","DOI":"10.1109\/TMM.2018.2866771","volume":"21","author":"D Hu","year":"2018","unstructured":"Hu D, Nie F, Li X (2018) Deep binary reconstruction for cross-modal hashing. IEEE Trans Multimedia 21(4):973","journal-title":"IEEE Trans Multimedia"},{"issue":"9","key":"6696_CR9","doi-asserted-by":"publisher","first-page":"2033","DOI":"10.1109\/TMM.2017.2703636","volume":"19","author":"X Zhu","year":"2017","unstructured":"Zhu X, Li X, Zhang S, Xu Z, Yu L, Wang C (2017) Graph pca hashing for similarity search. IEEE Trans Multimedia 19(9):2033","journal-title":"IEEE Trans Multimedia"},{"key":"6696_CR10","doi-asserted-by":"crossref","unstructured":"Shi Y, You X, Zheng F, Wang S, Peng Q (2019) Equally-guided discriminative hashing for cross-modal retrieval. In: Twenty-eighth international joint conference on artificial intelligence, pp 4767\u20134773","DOI":"10.24963\/ijcai.2019\/662"},{"issue":"1","key":"6696_CR11","doi-asserted-by":"publisher","first-page":"174","DOI":"10.1109\/TMM.2019.2922128","volume":"22","author":"J Zhang","year":"2020","unstructured":"Zhang J, Peng Y (2020) Multi-pathway generative adversarial hashing for unsupervised cross-modal retrieval. IEEE Trans. Multimedia 22(1):174","journal-title":"IEEE Trans. Multimedia"},{"issue":"9","key":"6696_CR12","doi-asserted-by":"publisher","first-page":"1404","DOI":"10.1109\/TMM.2015.2455415","volume":"17","author":"D Wang","year":"2015","unstructured":"Wang D, Cui P, Ou M, Zhu W (2015) Learning compact hash codes for multimodal representations using orthogonal deep structure. IEEE Trans. Multimedia 17(9):1404","journal-title":"IEEE Trans. Multimedia"},{"key":"6696_CR13","unstructured":"Liu W, Mu C, Kumar S, Chang S (2014) Discrete graph hashing. In: Advances in Neural Information Processing Systems, pp 3419\u20133427"},{"issue":"3","key":"6696_CR14","doi-asserted-by":"publisher","first-page":"571","DOI":"10.1109\/TMM.2016.2625747","volume":"19","author":"K Ding","year":"2017","unstructured":"Ding K, Fan B, Huo C, Xiang S, Pan C (2017) Cross-modal hashing via rank-order preserving. IEEE Trans Multimedia 19(3):571","journal-title":"IEEE Trans Multimedia"},{"key":"6696_CR15","doi-asserted-by":"crossref","unstructured":"Zhen L, Hu P, Wang X, Peng D (2019) Deep supervised cross-modal retrieval. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 10386\u201310395","DOI":"10.1109\/CVPR.2019.01064"},{"key":"6696_CR16","doi-asserted-by":"crossref","unstructured":"Cao Y, Long M, Wang J, Zhu H (2016) Correlation autoencoder hashing for supervised cross-modal search. In: Proceedings of the 2016 ACM on international conference on multimedia retrieval, pp 197\u2013204","DOI":"10.1145\/2911996.2912000"},{"key":"6696_CR17","doi-asserted-by":"crossref","unstructured":"Lin Z, Ding G, Hu M, Wang J (2015) Semantics-preserving hashing for cross-view retrieval. In: 2015 IEEE conference on computer vision and pattern recognition, pp 3864\u20133872","DOI":"10.1109\/CVPR.2015.7299011"},{"key":"6696_CR18","doi-asserted-by":"crossref","unstructured":"Mandal D, Chaudhury K, Biswas S (2017) Generalized semantic preserving hashing for n-label cross-modal retrieval. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2633\u20132641","DOI":"10.1109\/CVPR.2017.282"},{"key":"6696_CR19","doi-asserted-by":"crossref","unstructured":"Jiang Q, Li W (2017) Deep cross-modal hashing. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3232\u20133240","DOI":"10.1109\/CVPR.2017.348"},{"key":"6696_CR20","doi-asserted-by":"crossref","unstructured":"Yang E, Deng C, Liu W, Liu X, Tao D, Gao X (2017) Pairwise relationship guided deep hashing for cross-modal retrieval. In: Thirty-First AAAI conference on artificial intelligence, pp 1618\u20131625","DOI":"10.1609\/aaai.v31i1.10719"},{"key":"6696_CR21","doi-asserted-by":"crossref","unstructured":"Weng W, Wu J, Yang L, Liu L, Hu B (2019) Label-based deep semantic hashing for cross-modal retrieval. In: Neural Information Processing, pp 24\u201336","DOI":"10.1007\/978-3-030-36718-3_3"},{"key":"6696_CR22","doi-asserted-by":"crossref","unstructured":"Li C, Deng C, Li N, Liu W, Gao X, Tao D (2018) Self-supervised adversarial hashing networks for cross-modal retrieval. In: 2018 IEEE conference on computer vision and pattern recognition, pp 4242\u20134251","DOI":"10.1109\/CVPR.2018.00446"},{"key":"6696_CR23","doi-asserted-by":"crossref","unstructured":"Zhang X, Lai H, Feng J (2018) Attention-aware deep adversarial hashing for cross-modal retrieval. In: Proceedings of the European conference on computer vision, pp 591\u2013606","DOI":"10.1007\/978-3-030-01267-0_36"},{"key":"6696_CR24","doi-asserted-by":"crossref","unstructured":"Zhou J, Ding G, Guo Y (2014) Latent semantic sparse hashing for cross-modal similarity search. In: Proceedings of the 37th international ACM SIGIR conference on research & development in information retrieval, pp 415\u2013424","DOI":"10.1145\/2600428.2609610"},{"key":"6696_CR25","doi-asserted-by":"crossref","unstructured":"Song J, Yang Y, Yang Y, Huang Z, Shen HT (2013) Inter-media hashing for large-scale retrieval from heterogeneous data sources. In: Proceedings of the 2013 ACM SIGMOD international conference on management of data, pp 785\u2013796","DOI":"10.1145\/2463676.2465274"},{"key":"6696_CR26","doi-asserted-by":"crossref","unstructured":"Ding G, Guo Y, Zhou J (2014) Collective matrix factorization hashing for multimodal data. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2075\u20132082","DOI":"10.1109\/CVPR.2014.267"},{"key":"6696_CR27","doi-asserted-by":"crossref","unstructured":"Zhang D, Li W (2014) Large-scale supervised multimodal hashing with semantic correlation maximization. In: Twenty-Eighth AAAI Conference on Artificial Intelligence, pp 2177\u20132183","DOI":"10.1609\/aaai.v28i1.8995"},{"key":"6696_CR28","doi-asserted-by":"crossref","unstructured":"Lin Z, Ding G, Hu M, Wang J (2015) Semantics-preserving hashing for cross-view retrieval. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3864\u20133872","DOI":"10.1109\/CVPR.2015.7299011"},{"issue":"5","key":"6696_CR29","doi-asserted-by":"publisher","first-page":"2494","DOI":"10.1109\/TIP.2017.2676345","volume":"26","author":"X Xu","year":"2017","unstructured":"Xu X, Shen F, Yang Y, Shen HT, Li X (2017) Learning discriminative binary codes for large-scale cross-modal retrieval. IEEE Trans Image Process 26(5):2494","journal-title":"IEEE Trans Image Process"},{"key":"6696_CR30","doi-asserted-by":"crossref","unstructured":"Zhang J, Peng Y, Yuan M (2018) Unsupervised generative adversarial cross-modal hashing. In: 32nd AAAI conference on artificial intelligence, pp 539\u2013546","DOI":"10.1609\/aaai.v32i1.11263"},{"issue":"4","key":"6696_CR31","doi-asserted-by":"publisher","first-page":"1047","DOI":"10.1109\/TMM.2018.2869276","volume":"21","author":"M Yang","year":"2019","unstructured":"Yang M, Zhao W, Xu W, Feng Y, Zhao Z, Chen X, Lei K (2019) Multitask learning for cross-domain image captioning. IEEE Trans Multimedia 21(4):1047","journal-title":"IEEE Trans Multimedia"},{"issue":"6","key":"6696_CR32","doi-asserted-by":"publisher","first-page":"1367","DOI":"10.1109\/TPAMI.2017.2708709","volume":"40","author":"Q Wu","year":"2018","unstructured":"Wu Q, Shen C, Wang P, Dick A, Hengel A (2018) Image captioning and visual question answering based on attributes and external knowledge. IEEE Trans Pattern Anal Mach Intell 40(6):1367","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"issue":"2","key":"6696_CR33","doi-asserted-by":"publisher","first-page":"266","DOI":"10.1109\/TASLP.2017.2772846","volume":"26","author":"K Chen","year":"2018","unstructured":"Chen K, Zhao T, Yang M, Liu L, Tamura A, Wang R, Utiyama M, Sumita E (2018) A neural approach to source dependence based context model for statistical machine translation. IEEE\/ACM Trans Audio Speech Language Process 26(2):266","journal-title":"IEEE\/ACM Trans Audio Speech Language Process"},{"key":"6696_CR34","doi-asserted-by":"crossref","unstructured":"Yang Z, He X, Gao J, Deng L, Smola A (2016) Stacked attention networks for image question answering. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 21\u201329","DOI":"10.1109\/CVPR.2016.10"},{"issue":"11","key":"6696_CR35","doi-asserted-by":"publisher","first-page":"5585","DOI":"10.1109\/TIP.2018.2852503","volume":"27","author":"Y Peng","year":"2018","unstructured":"Peng Y, Qi J, Yuan Y (2018) Modality-specific cross-modal similarity measurement with recurrent attention network. IEEE Trans Image Process 27(11):5585","journal-title":"IEEE Trans Image Process"},{"key":"6696_CR36","doi-asserted-by":"crossref","unstructured":"Chatfield K, Simonyan K, Vedaldi A, Zisserman A (2014). Return of the devil in the details: Delving deep into convolutional nets, arXiv:1405.3531","DOI":"10.5244\/C.28.6"},{"key":"6696_CR37","doi-asserted-by":"crossref","unstructured":"He K, Jian S (2015) Convolutional neural networks at constrained time cost. In: Proceedings of IEEE conference on computer vision and pattern recognition, pp 5353\u20135360","DOI":"10.1109\/CVPR.2015.7299173"},{"key":"6696_CR38","doi-asserted-by":"crossref","unstructured":"Huiskes M, Lew M (2008) The MIR flickr retrieval evaluation. In: Proceedings of the 1st ACM international conference on multimedia information retrieval, pp 39\u201343","DOI":"10.1145\/1460096.1460104"},{"key":"6696_CR39","doi-asserted-by":"crossref","unstructured":"Chua T, Tang J, Hong R, Li H, Luo Z, Zheng Y (2009) NUS-WIDE: a real-world web image database from National University of Singapore. In: Proceedings of the ACM international conference on image and video retrieval, p 48","DOI":"10.1145\/1646396.1646452"},{"key":"6696_CR40","unstructured":"Liu W, Mu C, Kumar S, Chang SF (2014) Discrete graph hashing. In: Proceedings of the 27th international conference on neural information processing systems, pp 3419\u20133427"},{"key":"6696_CR41","unstructured":"Kumar S, Udupa R (2011) Learning hash functions for cross-view similarity search. In: Twenty-second international joint conference on artificial intelligence, pp 1360\u20131365"},{"key":"6696_CR42","unstructured":"Wang D, Gao X, Wang X, He L (2015) Semantic topic multimodal hashing for cross-media retrieval. In: Twenty-fourth international joint conference on artificial intelligence, pp 3890\u20133896"}],"container-title":["Neural Computing and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00521-021-06696-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00521-021-06696-y\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00521-021-06696-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2023,1,15]],"date-time":"2023-01-15T02:11:15Z","timestamp":1673748675000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00521-021-06696-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2021,11,12]]},"references-count":42,"journal-issue":{"issue":"7","published-print":{"date-parts":[[2022,4]]}},"alternative-id":["6696"],"URL":"https:\/\/doi.org\/10.1007\/s00521-021-06696-y","relation":{},"ISSN":["0941-0643","1433-3058"],"issn-type":[{"type":"print","value":"0941-0643"},{"type":"electronic","value":"1433-3058"}],"subject":[],"published":{"date-parts":[[2021,11,12]]},"assertion":[{"value":"30 March 2021","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 October 2021","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"12 November 2021","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}