{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,16]],"date-time":"2026-07-16T18:37:57Z","timestamp":1784227077809,"version":"3.55.0"},"reference-count":57,"publisher":"Springer Science and Business Media LLC","issue":"10","license":[{"start":{"date-parts":[[2025,5,7]],"date-time":"2025-05-07T00:00:00Z","timestamp":1746576000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,5,7]],"date-time":"2025-05-07T00:00:00Z","timestamp":1746576000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100017630","name":"Humanities and Social Sciences Youth Foundation, Ministry of Education","doi-asserted-by":"publisher","award":["20YJCZH172"],"award-info":[{"award-number":["20YJCZH172"]}],"id":[{"id":"10.13039\/501100017630","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100002858","name":"China Postdoctoral Science Foundation","doi-asserted-by":"publisher","award":["2019M651262"],"award-info":[{"award-number":["2019M651262"]}],"id":[{"id":"10.13039\/501100002858","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100010009","name":"Heilongjiang Provincial Postdoctoral Science Foundation","doi-asserted-by":"publisher","award":["LBH-Z19015"],"award-info":[{"award-number":["LBH-Z19015"]}],"id":[{"id":"10.13039\/501100010009","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Basic Research Support Program for Outstanding Young Teachers in Provincial Undergraduate Universities in Heilongjiang Province","award":["YQJH2023302"],"award-info":[{"award-number":["YQJH2023302"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2025,7]]},"DOI":"10.1007\/s10489-025-06577-0","type":"journal-article","created":{"date-parts":[[2025,5,7]],"date-time":"2025-05-07T02:03:51Z","timestamp":1746583431000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["TsAFN: A two-stage adaptive fusion network for multimodal sentiment analysis"],"prefix":"10.1007","volume":"55","author":[{"given":"Jiaqi","family":"Liu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-9867-9765","authenticated-orcid":false,"given":"Yong","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jing","family":"Yang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xu","family":"Yu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Meng","family":"Zhao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,5,7]]},"reference":[{"key":"6577_CR1","doi-asserted-by":"crossref","unstructured":"Kaur R, Kautish S (2022) Multimodal sentiment analysis: A survey and comparison. Res Anthol Implementing Sentiment Anal Across Multiple Disc 1846\u20131870","DOI":"10.4018\/978-1-6684-6303-1.ch098"},{"key":"6577_CR2","doi-asserted-by":"crossref","unstructured":"Das R, Singh TD (2023) Multimodal sentiment analysis: A survey of methods, trends and challenges. ACM Comput Surv","DOI":"10.1145\/3586075"},{"key":"6577_CR3","doi-asserted-by":"publisher","first-page":"118290","DOI":"10.1016\/j.eswa.2022.118290","volume":"209","author":"R Catelli","year":"2022","unstructured":"Catelli R, Fujita H, De Pietro G, Esposito M (2022) Deceptive reviews and sentiment polarity: Effective link by exploiting bert. Expert Syst Appl 209:118290","journal-title":"Expert Syst Appl"},{"key":"6577_CR4","doi-asserted-by":"crossref","unstructured":"Ghorbanali A, Sohrabi MK (2023) A comprehensive survey on deep learning-based approaches for multimodal sentiment analysis. Artif Intell Rev 1\u201334","DOI":"10.1007\/s10462-023-10555-8"},{"key":"6577_CR5","doi-asserted-by":"crossref","unstructured":"Poria S, Cambria E, Gelbukh A (2015) Deep convolutional neural network textual features and multiple kernel learning for utterance-level multimodal sentiment analysis, 2539\u20132544","DOI":"10.18653\/v1\/D15-1303"},{"key":"6577_CR6","doi-asserted-by":"crossref","unstructured":"Siddiquie B, Chisholm D, Divakaran A (2015) Exploiting multimodal affect and semantics to identify politically persuasive web videos, 203\u2013210","DOI":"10.1145\/2818346.2820732"},{"key":"6577_CR7","doi-asserted-by":"publisher","first-page":"50","DOI":"10.1016\/j.neucom.2015.01.095","volume":"174","author":"S Poria","year":"2016","unstructured":"Poria S, Cambria E, Howard N, Huang G-B, Hussain A (2016) Fusing audio, visual and textual clues for sentiment analysis from multimodal content. Neurocomputing 174:50\u201359","journal-title":"Neurocomputing"},{"key":"6577_CR8","doi-asserted-by":"crossref","unstructured":"Cambria E, Hazarika D, Poria S, Hussain A, Subramanyam R (2018) Benchmarking multimodal sentiment analysis, 166\u2013179 (organizationSpringer)","DOI":"10.1007\/978-3-319-77116-8_13"},{"key":"6577_CR9","doi-asserted-by":"publisher","first-page":"124","DOI":"10.1016\/j.knosys.2018.07.041","volume":"161","author":"N Majumder","year":"2018","unstructured":"Majumder N, Hazarika D, Gelbukh A, Cambria E, Poria S (2018) Multimodal sentiment analysis using hierarchical fusion with context modeling. Knowl-based Syst 161:124\u2013133","journal-title":"Knowl-based Syst"},{"key":"6577_CR10","doi-asserted-by":"crossref","unstructured":"Yang D, Huang S, Kuang H, Du Y, Zhang L (2022) Disentangled representation learning for multimodal emotion recognition, 1642\u20131651","DOI":"10.1145\/3503161.3547754"},{"key":"6577_CR11","doi-asserted-by":"crossref","unstructured":"Mai S, Zeng Y, Zheng S, Hu H (2022) Hybrid contrastive learning of tri-modal representation for multimodal sentiment analysis. IEEE Trans Affective Comput","DOI":"10.1109\/TAFFC.2022.3172360"},{"key":"6577_CR12","doi-asserted-by":"publisher","first-page":"37","DOI":"10.1016\/j.inffus.2022.11.022","volume":"92","author":"K Kim","year":"2023","unstructured":"Kim K, Park S (2023) Aobert: All-modalities-in-one bert for multimodal sentiment analysis. Inf Fusion 92:37\u201345","journal-title":"Inf Fusion"},{"key":"6577_CR13","doi-asserted-by":"publisher","first-page":"992","DOI":"10.1109\/LSP.2021.3078074","volume":"28","author":"J He","year":"2021","unstructured":"He J, Mai S, Hu H (2021) A unimodal reinforced transformer with time squeeze fusion for multimodal sentiment analysis. IEEE Signal Process Lett 28:992\u2013996","journal-title":"IEEE Signal Process Lett"},{"key":"6577_CR14","doi-asserted-by":"crossref","unstructured":"Yu W, Xu H, Yuan Z, Wu J (2021) Learning modality-specific representations with self-supervised multi-task learning for multimodal sentiment analysis 35:10790\u201310797","DOI":"10.1609\/aaai.v35i12.17289"},{"key":"6577_CR15","doi-asserted-by":"publisher","first-page":"460","DOI":"10.1109\/JPROC.2014.2306253","volume":"102","author":"AH Sayed","year":"2014","unstructured":"Sayed AH (2014) Adaptive networks. Proceed IEEE 102:460\u2013497","journal-title":"Proceed IEEE"},{"key":"6577_CR16","doi-asserted-by":"publisher","first-page":"118246","DOI":"10.1016\/j.eswa.2022.118246","volume":"209","author":"R Catelli","year":"2022","unstructured":"Catelli R et al (2022) Cross lingual transfer learning for sentiment analysis of italian tripadvisor reviews. Expert Syst Appl 209:118246","journal-title":"Expert Syst Appl"},{"key":"6577_CR17","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3545572","volume":"19","author":"S Jabeen","year":"2023","unstructured":"Jabeen S et al (2023) A review on methods and applications in multimodal deep learning. ACM Trans Multimed Comput, Commun Appl 19:1\u201341","journal-title":"ACM Trans Multimed Comput, Commun Appl"},{"key":"6577_CR18","doi-asserted-by":"crossref","unstructured":"Zhu T et\u00a0al (2022) Multimodal sentiment analysis with image-text interaction network. IEEE Trans Multimed","DOI":"10.1109\/TMM.2022.3160060"},{"key":"6577_CR19","doi-asserted-by":"crossref","unstructured":"Yu J, Chen K, Xia R (2022) Hierarchical interactive multimodal transformer for aspect-based multimodal sentiment analysis. IEEE Trans Affective Comput","DOI":"10.1109\/TAFFC.2022.3171091"},{"key":"6577_CR20","doi-asserted-by":"crossref","unstructured":"Han W, Chen H, Poria S (2021) Improving multimodal fusion with hierarchical mutual information maximization for multimodal sentiment analysis. arXiv:2109.00412","DOI":"10.18653\/v1\/2021.emnlp-main.723"},{"key":"6577_CR21","doi-asserted-by":"publisher","first-page":"103193","DOI":"10.1016\/j.ipm.2022.103193","volume":"60","author":"D Chen","year":"2023","unstructured":"Chen D, Su W, Wu P, Hua B (2023) Joint multimodal sentiment analysis based on information relevance. Inf Process Manag 60:103193","journal-title":"Inf Process Manag"},{"key":"6577_CR22","doi-asserted-by":"crossref","unstructured":"Xiao L, Wu X, Wu W, Yang J, He L (2022) Multi-channel attentive graph convolutional network with sentiment fusion for multimodal sentiment analysis, 4578\u20134582 (organizationIEEE)","DOI":"10.1109\/ICASSP43922.2022.9747542"},{"key":"6577_CR23","doi-asserted-by":"publisher","first-page":"103229","DOI":"10.1016\/j.ipm.2022.103229","volume":"60","author":"H Lin","year":"2023","unstructured":"Lin H et al (2023) Ps-mixer: A polar-vector and strength-vector mixer model for multimodal sentiment analysis. Inf Process Manag 60:103229","journal-title":"Inf Process Manag"},{"key":"6577_CR24","doi-asserted-by":"crossref","unstructured":"Sun H, Wang H, Liu J, Chen Y-W, Lin L (2022) Cubemlp: An mlp-based model for multimodal sentiment analysis and depression estimation, 3722\u20133729","DOI":"10.1145\/3503161.3548025"},{"key":"6577_CR25","first-page":"5105","volume":"35","author":"X Xue","year":"2022","unstructured":"Xue X, Zhang C, Niu Z, Wu X (2022) Multi-level attention map network for multimodal sentiment analysis. IEEE Trans Knowl Data Eng 35:5105\u20135118","journal-title":"IEEE Trans Knowl Data Eng"},{"key":"6577_CR26","doi-asserted-by":"crossref","unstructured":"Fu Z et\u00a0al (2022) Nhfnet: A non-homogeneous fusion network for multimodal sentiment analysis, 1\u20136 (organizationIEEE)","DOI":"10.1109\/ICME52920.2022.9859836"},{"key":"6577_CR27","doi-asserted-by":"publisher","first-page":"108107","DOI":"10.1016\/j.knosys.2021.108107","volume":"240","author":"Y Du","year":"2022","unstructured":"Du Y, Liu Y, Peng Z, Jin X (2022) Gated attention fusion network for multimodal sentiment classification. Knowl-Based Syst 240:108107","journal-title":"Knowl-Based Syst"},{"key":"6577_CR28","doi-asserted-by":"crossref","unstructured":"Kumar P, Khokher V, Gupta Y, Raman B (2021) Hybrid fusion based approach for multimodal emotion recognition with insufficient labeled data, 314\u2013318 (organizationIEEE)","DOI":"10.1109\/ICIP42928.2021.9506714"},{"key":"6577_CR29","doi-asserted-by":"crossref","unstructured":"Paraskevopoulos G, Georgiou E, Potamianos A (2022) Mmlatch: Bottom-up top-down fusion for multimodal sentiment analysis, 4573\u20134577 (organizationIEEE)","DOI":"10.1109\/ICASSP43922.2022.9746418"},{"key":"6577_CR30","doi-asserted-by":"crossref","unstructured":"Manzoor MA et\u00a0al (2023) Multimodality representation learning: A survey on evolution, pretraining and its applications. arXiv:2302.00389","DOI":"10.1145\/3617833"},{"key":"6577_CR31","doi-asserted-by":"crossref","unstructured":"Chandrasekaran G, Nguyen TN, Hemanth DJ (2021) Multimodal sentimental analysis for social media applications: A comprehensive review. Wiley Interdisciplinary Reviews: Data Mining and Knowledge Discovery 11:e1415","DOI":"10.1002\/widm.1415"},{"key":"6577_CR32","doi-asserted-by":"crossref","unstructured":"Im J, Kim M, Lee H, Cho H, Chung S (2021) Self-supervised multimodal opinion summarization. arXiv:2105.13135","DOI":"10.18653\/v1\/2021.acl-long.33"},{"key":"6577_CR33","doi-asserted-by":"publisher","first-page":"1093","DOI":"10.3390\/app12031093","volume":"12","author":"J Wang","year":"2022","unstructured":"Wang J, Mao H, Li H (2022) Fmfn: Fine-grained multimodal fusion networks for fake news detection. Appl Sci 12:1093","journal-title":"Appl Sci"},{"key":"6577_CR34","doi-asserted-by":"crossref","unstructured":"Zhang X et\u00a0al (2021) Learning robust patient representations from multi-modal electronic health records: a supervised deep learning approach, 585\u2013593 (organizationSIAM)","DOI":"10.1137\/1.9781611976700.66"},{"key":"6577_CR35","doi-asserted-by":"publisher","first-page":"856","DOI":"10.1109\/TEVC.2021.3066285","volume":"25","author":"F Huang","year":"2021","unstructured":"Huang F, Jolfaei A, Bashir AK (2021) Robust multimodal representation learning with evolutionary adversarial attention networks. IEEE Trans Evolutionary Comput 25:856\u2013868","journal-title":"IEEE Trans Evolutionary Comput"},{"key":"6577_CR36","doi-asserted-by":"publisher","first-page":"208","DOI":"10.1016\/j.ins.2023.01.116","volume":"628","author":"J Wang","year":"2023","unstructured":"Wang J, Wang S, Lin M, Xu Z, Guo W (2023) Learning speaker-independent multimodal representation for sentiment analysis. Inf Sci 628:208\u2013225","journal-title":"Inf Sci"},{"key":"6577_CR37","doi-asserted-by":"crossref","unstructured":"Guo X, Kong W-KA, Kot AC (2022) Deep multimodal sequence fusion by regularized expressive representation distillation. IEEE Trans Multimed","DOI":"10.1109\/TMM.2022.3142448"},{"key":"6577_CR38","doi-asserted-by":"crossref","unstructured":"Yu L et\u00a0al (2022) Commercemm: Large-scale commerce multimodal representation learning with omni retrieval, 4433\u20134442","DOI":"10.1145\/3534678.3539151"},{"key":"6577_CR39","doi-asserted-by":"crossref","unstructured":"Yang Z et\u00a0al (2023) i-code: An integrative and composable multimodal learning framework, Vol.\u00a037, 10880\u201310890","DOI":"10.1609\/aaai.v37i9.26290"},{"key":"6577_CR40","doi-asserted-by":"publisher","first-page":"110125","DOI":"10.1016\/j.knosys.2022.110125","volume":"260","author":"B Huang","year":"2023","unstructured":"Huang B et al (2023) Crf-gcn: An effective syntactic dependency model for aspect-level sentiment analysis. Knowl-Based Syst 260:110125","journal-title":"Knowl-Based Syst"},{"key":"6577_CR41","doi-asserted-by":"publisher","first-page":"297","DOI":"10.1016\/j.neucom.2021.04.141","volume":"485","author":"D Liu","year":"2022","unstructured":"Liu D, Kong H, Luo X, Liu W, Subramaniam R (2022) Bringing ai to edge: From deep learning\u2019s perspective. Neurocomputing 485:297\u2013320","journal-title":"Neurocomputing"},{"key":"6577_CR42","doi-asserted-by":"publisher","first-page":"1574","DOI":"10.1109\/LSP.2022.3179946","volume":"29","author":"Y Gao","year":"2022","unstructured":"Gao Y, Fu X, Ouyang T, Wang Y (2022) Eeg-gcn: spatio-temporal and self-adaptive graph convolutional networks for single and multi-view eeg-based emotion recognition. IEEE Signal Process Lett 29:1574\u20131578","journal-title":"IEEE Signal Process Lett"},{"key":"6577_CR43","doi-asserted-by":"publisher","first-page":"455","DOI":"10.1007\/s11263-021-01556-7","volume":"130","author":"D Ruan","year":"2022","unstructured":"Ruan D et al (2022) Adaptive deep disturbance-disentangled learning for facial expression recognition. Int J Comput Vision 130:455\u2013477","journal-title":"Int J Comput Vision"},{"key":"6577_CR44","doi-asserted-by":"crossref","unstructured":"Huijuan Z, Ning Y, Ruchuan W (2023) Improved cross-corpus speech emotion recognition using deep local domain adaptation. Chinese J Electron 32:1\u20137","DOI":"10.23919\/cje.2021.00.196"},{"key":"6577_CR45","doi-asserted-by":"publisher","first-page":"103235","DOI":"10.1016\/j.bspc.2021.103235","volume":"71","author":"M Yan","year":"2022","unstructured":"Yan M et al (2022) Emotion classification with multichannel physiological signals using hybrid feature and adaptive decision fusion. Biomed Signal Process Control 71:103235","journal-title":"Biomed Signal Process Control"},{"key":"6577_CR46","doi-asserted-by":"publisher","first-page":"110149","DOI":"10.1016\/j.knosys.2022.110149","volume":"260","author":"C Arumugam","year":"2023","unstructured":"Arumugam C, Nallaperumal K (2023) Eiaasg: Emotional intensive adaptive aspect-specific gcn for sentiment classification. Knowl-Based Syst 260:110149","journal-title":"Knowl-Based Syst"},{"key":"6577_CR47","unstructured":"Devlin J, Chang M-W, Lee K, Toutanova K (2018) Bert: Pre-training of deep bidirectional transformers for language understanding. arXiv:1810.04805"},{"key":"6577_CR48","doi-asserted-by":"crossref","unstructured":"McFee B et\u00a0al (2015) librosa: Audio and music signal analysis in python, Vol.\u00a08, 18\u201325","DOI":"10.25080\/Majora-7b98e3ed-003"},{"key":"6577_CR49","doi-asserted-by":"crossref","unstructured":"Schroff F, Kalenichenko D, Philbin J (2015) Facenet: A unified embedding for face recognition and clustering, 815\u2013823","DOI":"10.1109\/CVPR.2015.7298682"},{"key":"6577_CR50","unstructured":"Zadeh A, Zellers R, Pincus E, Morency L-P (2016) Mosi: multimodal corpus of sentiment intensity and subjectivity analysis in online opinion videos. arXiv:1606.06259"},{"key":"6577_CR51","unstructured":"Zadeh AB, Liang PP, Poria S, Cambria E, Morency L-P (2018) Multimodal language analysis in the wild: Cmu-mosei dataset and interpretable dynamic fusion graph, 2236\u20132246"},{"key":"6577_CR52","doi-asserted-by":"crossref","unstructured":"Yu W et\u00a0al (2020) Ch-sims: A chinese multimodal sentiment analysis dataset with fine-grained annotation of modality, 3718\u20133727","DOI":"10.18653\/v1\/2020.acl-main.343"},{"key":"6577_CR53","doi-asserted-by":"crossref","unstructured":"Tsai Y-HH et\u00a0al (2019) Multimodal transformer for unaligned multimodal language sequences","DOI":"10.18653\/v1\/P19-1656"},{"key":"6577_CR54","doi-asserted-by":"crossref","unstructured":"Hazarika D, Zimmermann R, Poria S (2020) Misa: Modality-invariant and-specific representations for multimodal sentiment analysis, 1122\u20131131","DOI":"10.1145\/3394171.3413678"},{"key":"6577_CR55","doi-asserted-by":"crossref","unstructured":"Rahman W et\u00a0al (2020) Integrating multimodal information in large pretrained transformers, Vol. 2020, 2359 (organizationNIH Public Access)","DOI":"10.18653\/v1\/2020.acl-main.214"},{"key":"6577_CR56","doi-asserted-by":"publisher","first-page":"296","DOI":"10.1016\/j.inffus.2022.07.006","volume":"88","author":"F Zhang","year":"2022","unstructured":"Zhang F et al (2022) Deep emotional arousal network for multimodal sentiment analysis and emotion recognition. Inf Fusion 88:296\u2013304","journal-title":"Inf Fusion"},{"key":"6577_CR57","doi-asserted-by":"publisher","first-page":"107676","DOI":"10.1016\/j.knosys.2021.107676","volume":"235","author":"T Wu","year":"2022","unstructured":"Wu T et al (2022) Video sentiment analysis with bimodal information-augmented multi-head attention. Knowl-Based Syst 235:107676","journal-title":"Knowl-Based Syst"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-025-06577-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-025-06577-0\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-025-06577-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,9,19]],"date-time":"2025-09-19T13:59:04Z","timestamp":1758290344000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-025-06577-0"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,5,7]]},"references-count":57,"journal-issue":{"issue":"10","published-print":{"date-parts":[[2025,7]]}},"alternative-id":["6577"],"URL":"https:\/\/doi.org\/10.1007\/s10489-025-06577-0","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,5,7]]},"assertion":[{"value":"16 April 2025","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"7 May 2025","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that there is no conflict of interests regarding the publication of the article.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of Interest"}},{"value":"The data utilized in this study were sourced from publicly available datasets, which do not involve any ethical or privacy concerns. These datasets were collected and made accessible by , and their use is governed by the terms and conditions set forth by the data provider.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical and Informed Consent for Data Used"}}],"article-number":"725"}}