{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,7]],"date-time":"2026-07-07T15:37:40Z","timestamp":1783438660283,"version":"3.54.6"},"reference-count":35,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2024,12,19]],"date-time":"2024-12-19T00:00:00Z","timestamp":1734566400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,12,19]],"date-time":"2024-12-19T00:00:00Z","timestamp":1734566400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Natural Science Foundation","award":["61702462 and 61906175"],"award-info":[{"award-number":["61702462 and 61906175"]}]},{"DOI":"10.13039\/501100006683","name":"XJTLU","doi-asserted-by":"crossref","award":["RDF-21-02-008"],"award-info":[{"award-number":["RDF-21-02-008"]}],"id":[{"id":"10.13039\/501100006683","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100017700","name":"Henan Provincial Science and Technology Research Project","doi-asserted-by":"publisher","award":["232102211006, 232102210044, 232102211017, 232102210055 and 222102210214"],"award-info":[{"award-number":["232102211006, 232102210044, 232102211017, 232102210055 and 222102210214"]}],"id":[{"id":"10.13039\/501100017700","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Jiangsu Double-Innovation Plan","award":["JSSCBS20230474"],"award-info":[{"award-number":["JSSCBS20230474"]}]},{"DOI":"10.13039\/501100012502","name":"Collaborative Innovation Center for Water Treatment Technology and Materials","doi-asserted-by":"publisher","award":["23XNKJTD0205"],"award-info":[{"award-number":["23XNKJTD0205"]}],"id":[{"id":"10.13039\/501100012502","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Undergraduate Universities Smart Teaching Special Research Project of Henan Province","award":["No. 489-29"],"award-info":[{"award-number":["No. 489-29"]}]},{"name":"Doctor Natural Science Foundation of Zhengzhou University of Light Industry","award":["2021BSJJ025 and 2022BSJJZK13"],"award-info":[{"award-number":["2021BSJJ025 and 2022BSJJZK13"]}]},{"DOI":"10.13039\/501100006407","name":"Natural Science Foundation of Henan","doi-asserted-by":"crossref","award":["242300421220"],"award-info":[{"award-number":["242300421220"]}],"id":[{"id":"10.13039\/501100006407","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2025,2]]},"DOI":"10.1007\/s10489-024-06150-1","type":"journal-article","created":{"date-parts":[[2024,12,19]],"date-time":"2024-12-19T09:57:12Z","timestamp":1734602232000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":7,"title":["Text-dominant multimodal perception network for sentiment analysis based on cross-modal semantic enhancements"],"prefix":"10.1007","volume":"55","author":[{"given":"Zuhe","family":"Li","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Panbo","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-6877-3937","authenticated-orcid":false,"given":"Yushan","family":"Pan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jun","family":"Yu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Weihua","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Haoran","family":"Chen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yiming","family":"Luo","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Hao","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,12,19]]},"reference":[{"key":"6150_CR1","doi-asserted-by":"publisher","first-page":"306","DOI":"10.1016\/j.inffus.2023.02.028","volume":"95","author":"L Zhu","year":"2023","unstructured":"Zhu L, Zhu Z, Zhang C, Xu Y, Kong X (2023) Multimodal sentiment analysis based on fusion methods: A survey. Inf Fusion 95:306\u2013325. https:\/\/doi.org\/10.1016\/j.inffus.2023.02.028","journal-title":"Inf Fusion"},{"key":"6150_CR2","doi-asserted-by":"publisher","first-page":"112011","DOI":"10.1016\/j.asoc.2024.112011","volume":"164","author":"Z Liu","year":"2024","unstructured":"Liu Z, Yang T, Chen W, Chen J, Li Q, Zhang J (2024) Sentiment analysis of social media comments based on multimodal attention fusion network. Appl Soft Comput 164:112011. https:\/\/doi.org\/10.1016\/j.asoc.2024.112011","journal-title":"Appl Soft Comput"},{"key":"6150_CR3","doi-asserted-by":"publisher","first-page":"111553","DOI":"10.1016\/j.asoc.2024.111553","volume":"157","author":"AG Aruna","year":"2024","unstructured":"Aruna AG, Vetriselvi V (2024) Sentiment analysis on a low-resource language dataset using multimodal representation learning and cross-lingual transfer learning. Appl Soft Comput 157:111553. https:\/\/doi.org\/10.1016\/j.asoc.2024.111553","journal-title":"Appl Soft Comput"},{"key":"6150_CR4","doi-asserted-by":"publisher","first-page":"424","DOI":"10.1016\/j.inffus.2022.09.025","volume":"91","author":"A Gandhi","year":"2023","unstructured":"Gandhi A, Adhvaryu K, Poria S, Cambria E, Hussain A (2023) Multimodal sentiment analysis: A systematic review of history, datasets, multimodal fusion methods, applications, challenges and future directions. Inf Fusion 91:424\u2013444. https:\/\/doi.org\/10.1016\/j.inffus.2022.09.025","journal-title":"Inf Fusion"},{"key":"6150_CR5","doi-asserted-by":"crossref","unstructured":"Zadeh A, Chen M, Cambria E, Poria S, Morency L-P (2017) Tensor fusion network for multimodal sentiment analysis. In: EMNLP 2017 - conference on empirical methods in natural language processing, Proceedings, Copenhagen, Denmark, pp 1103\u20131114","DOI":"10.18653\/v1\/D17-1115"},{"key":"6150_CR6","doi-asserted-by":"crossref","unstructured":"Liu Z, Shen Y, Lakshminarasimhan VB, Liang PP, Zadeh A, Morency L-P (2018) Efficient low-rank multimodal fusion with modality-specific factors. In: Gurevych I, Miyao Y (eds) Proceedings of the 56th annual meeting of the Association for Computational Linguistics (Acl), vol 1, pp 2247\u20132256","DOI":"10.18653\/v1\/P18-1209"},{"key":"6150_CR7","doi-asserted-by":"publisher","first-page":"101891","DOI":"10.1016\/j.inffus.2023.101891","volume":"99","author":"Z Li","year":"2023","unstructured":"Li Z, Guo Q, Pan Y, Ding W, Yu J, Zhang Y, Liu W, Chen H, Wang H, Xie Y (2023) Multi-level correlation mining framework with self-supervised label generation for multimodal sentiment analysis. Inf Fusion 99:101891. https:\/\/doi.org\/10.1016\/j.inffus.2023.101891","journal-title":"Inf Fusion"},{"key":"6150_CR8","doi-asserted-by":"publisher","first-page":"110502","DOI":"10.1016\/j.knosys.2023.110502","volume":"269","author":"C Huang","year":"2023","unstructured":"Huang C, Zhang J, Wu X, Wang Y, Li M, Huang X (2023) Tefna: Text-centered fusion network with crossmodal attention for multimodal sentiment analysis. Knowl-Based Syst 269:110502. https:\/\/doi.org\/10.1016\/j.knosys.2023.110502","journal-title":"Knowl-Based Syst"},{"key":"6150_CR9","doi-asserted-by":"publisher","first-page":"529","DOI":"10.1109\/TMM.2023.3267882","volume":"26","author":"Z Yuan","year":"2024","unstructured":"Yuan Z, Liu Y, Xu H, Gao K (2024) Noise imitation based adversarial training for robust multimodal sentiment analysis. IEEE Trans Multimed 26:529\u2013539. https:\/\/doi.org\/10.1109\/TMM.2023.3267882","journal-title":"IEEE Trans Multimed"},{"key":"6150_CR10","doi-asserted-by":"publisher","unstructured":"Liu W, Li W, Ruan Y-P, Shu Y, Chen J, Li Y, Yu C, Zhang Y, Guan J, Zhou S (2024) Weakly correlated multimodal sentiment analysis: New dataset and topic-oriented model. IEEE Trans Affect Comput 1\u201313. https:\/\/doi.org\/10.1109\/TAFFC.2024.3396144","DOI":"10.1109\/TAFFC.2024.3396144"},{"key":"6150_CR11","doi-asserted-by":"publisher","first-page":"47","DOI":"10.1016\/j.neucom.2021.05.040","volume":"455","author":"J Zhou","year":"2021","unstructured":"Zhou J, Zhao J, Huang JX, Hu QV, He L (2021) Masad: A large-scale dataset for multimodal aspect-based sentiment analysis. Neurocomputing 455:47\u201358. https:\/\/doi.org\/10.1016\/j.neucom.2021.05.040","journal-title":"Neurocomputing"},{"key":"6150_CR12","doi-asserted-by":"publisher","first-page":"98","DOI":"10.1016\/j.inffus.2017.02.003","volume":"37","author":"S Poria","year":"2017","unstructured":"Poria S, Cambria E, Bajpai R, Hussain A (2017) A review of affective computing: From unimodal analysis to multimodal fusion. Inf Fusion 37:98\u2013125. https:\/\/doi.org\/10.1016\/j.inffus.2017.02.003","journal-title":"Inf Fusion"},{"key":"6150_CR13","doi-asserted-by":"publisher","unstructured":"Mai S, Hu H, Xing S (2019) Divide, conquer and combine: Hierarchical feature fusion network with local and global perspectives for multimodal affective computing. In: Korhonen A, Traum D, M\u00e0rquez L (eds) Proceedings of the 57th annual meeting of the association for computational linguistics, pp 481\u2013492. Association for Computational Linguistics, Florence, Italy. https:\/\/doi.org\/10.18653\/v1\/P19-1046","DOI":"10.18653\/v1\/P19-1046"},{"key":"6150_CR14","doi-asserted-by":"publisher","unstructured":"Hazarika D, Zimmermann R, Poria S (2020) Misa: Modality-invariant and -specific representations for multimodal sentiment analysis. In: Proceedings of the 28th ACM international conference on multimedia, pp 1122\u20131131. Association for Computing Machinery, New York, NY, USA. https:\/\/doi.org\/10.1145\/3394171.3413678","DOI":"10.1145\/3394171.3413678"},{"key":"6150_CR15","doi-asserted-by":"crossref","unstructured":"Tsai Y-HH, Bai S, Liang PP, Kolter JZ, Morency L-P, Salakhutdinov R (2019) Multimodal transformer for unaligned multimodal language sequences. In: Korhonen A, Traum D, Marquez L (eds) 57th Annual Meeting of the Association for Computational Linguistics (Acl 2019), pp 6558\u20136569","DOI":"10.18653\/v1\/P19-1656"},{"issue":"4","key":"6150_CR16","doi-asserted-by":"publisher","first-page":"1966","DOI":"10.1109\/TCSVT.2022.3218018","volume":"33","author":"J Tang","year":"2023","unstructured":"Tang J, Liu D, Jin X, Peng Y, Zhao Q, Ding Y, Kong W (2023) Bafn: Bi-direction attention based fusion network for multimodal sentiment analysis. IEEE Trans Circ Syst Video Technol 33(4):1966\u20131978. https:\/\/doi.org\/10.1109\/TCSVT.2022.3218018","journal-title":"IEEE Trans Circ Syst Video Technol"},{"key":"6150_CR17","doi-asserted-by":"crossref","unstructured":"Zadeh A, Liang PP, Poria S, Vij P, Cambria E, Morency L-P (2018) Multi-attention recurrent network human communication comprehension. In: Thirty-second AAAI conference on artificial intelligence \/ thirtieth innovative applications of artificial intelligence conference \/ eighth AAAI symposium on educational advances in artificial intelligence, pp 5642\u20135649","DOI":"10.1609\/aaai.v32i1.12024"},{"key":"6150_CR18","doi-asserted-by":"publisher","unstructured":"Wang Z, Wan Z, Wan X (2020) Transmodality: An end2end fusion method with transformer for multimodal sentiment analysis. In: Proceedings of the web conference 2020, pp 2514\u20132520. Association for Computing Machinery, New York, NY, USA. https:\/\/doi.org\/10.1145\/3366423.3380000","DOI":"10.1145\/3366423.3380000"},{"key":"6150_CR19","doi-asserted-by":"publisher","first-page":"1476","DOI":"10.1109\/TASLP.2023.3263801","volume":"31","author":"C Chen","year":"2023","unstructured":"Chen C, Hong H, Guo J, Song B (2023) Inter-intra modal representation augmentation with trimodal collaborative disentanglement network for multimodal sentiment analysis. IEEE\/ACM Trans Audio Speech Lang Process 31:1476\u20131488. https:\/\/doi.org\/10.1109\/TASLP.2023.3263801","journal-title":"IEEE\/ACM Trans Audio Speech Lang Process"},{"key":"6150_CR20","doi-asserted-by":"crossref","unstructured":"Wu Y, Lin Z, Zhao Y, Qin B, Zhu L-N (2021) A text-centered shared-private framework via cross-modal prediction for multimodal sentiment analysis. In: Findings of the Association for Computational Linguistics: ACL-IJCNLP 2021, pp 4730\u20134738","DOI":"10.18653\/v1\/2021.findings-acl.417"},{"key":"6150_CR21","doi-asserted-by":"publisher","unstructured":"Han W, Chen H, Gelbukh A, Zadeh A, Morency L-P, Poria S (2021) Bi-bimodal modality fusion for correlation-controlled multimodal sentiment analysis. In: Proceedings of the 2021 international conference on multimodal interaction, pp 6\u201315. https:\/\/doi.org\/10.1145\/3462244.3479919","DOI":"10.1145\/3462244.3479919"},{"key":"6150_CR22","unstructured":"Devlin J, Chang M-W, Lee K, Toutanova K (2019) Bert: Pre-training of deep bidirectional transformers for language understanding. In: 2019 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (Naacl Hlt 2019), vol 1, pp 4171\u20134186"},{"issue":"8","key":"6150_CR23","doi-asserted-by":"publisher","first-page":"1735","DOI":"10.1162\/neco.1997.9.8.1735","volume":"9","author":"S Hochreiter","year":"1997","unstructured":"Hochreiter S, Schmidhuber J (1997) Long Short-Term Memory. Neural Comput 9(8):1735\u20131780. https:\/\/doi.org\/10.1162\/neco.1997.9.8.1735","journal-title":"Neural Comput"},{"key":"6150_CR24","unstructured":"Khosla P, Teterwak P, Wang C, Sarna A, Tian Y, Isola P, Maschinot A, Liu C, Krishnan D (2020) Supervised contrastive learning. In: Larochelle H, Ranzato M, Hadsell R, Balcan MF, Lin H (eds) Advances in neural information processing systems, vol 33, pp 18661\u201318673"},{"key":"6150_CR25","unstructured":"Zadeh A, Zellers R, Pincus E, Morency L (2016) MOSI: multimodal corpus of sentiment intensity and subjectivity analysis in online opinion videos"},{"key":"6150_CR26","doi-asserted-by":"publisher","unstructured":"Bagher\u00a0Zadeh A, Liang PP, Poria S, Cambria E, Morency L-P (2018) Multimodal language analysis in the wild: CMU-MOSEI dataset and interpretable dynamic fusion graph. In: Gurevych I, Miyao Y (eds) Proceedings of the 56th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers), pp 2236\u20132246. Association for Computational Linguistics, Melbourne, Australia. https:\/\/doi.org\/10.18653\/v1\/P18-1208","DOI":"10.18653\/v1\/P18-1208"},{"key":"6150_CR27","doi-asserted-by":"crossref","unstructured":"Sun Z, Sarma PK, Sethares WA, Liang Y (2020) Learning relationships between text, audio, and video via deep canonical correlation for multimodal language analysis. In: Thirty-fourth AAAI conference on artificial intelligence, the thirty-second innovative applications of artificial intelligence conference and the tenth AAAI symposium on educational advances in artificial intelligence, vol 34, pp 8992\u20138999","DOI":"10.1609\/aaai.v34i05.6431"},{"key":"6150_CR28","doi-asserted-by":"publisher","unstructured":"Yu W, Xu H, Yuan Z, Wu J (2021) Learning modality-specific representations with self-supervised multi-task learning for multimodal sentiment analysis. Proceedings of the AAAI conference on artificial intelligence, vol 35, no 12, pp 10790\u201310797. https:\/\/doi.org\/10.1609\/aaai.v35i12.17289","DOI":"10.1609\/aaai.v35i12.17289"},{"key":"6150_CR29","doi-asserted-by":"publisher","first-page":"130","DOI":"10.1016\/j.neucom.2021.09.041","volume":"467","author":"B Yang","year":"2022","unstructured":"Yang B, Shao B, Wu L, Lin X (2022) Multimodal sentiment analysis with unidirectional modality translation. Neurocomputing 467:130\u2013137. https:\/\/doi.org\/10.1016\/j.neucom.2021.09.041","journal-title":"Neurocomputing"},{"key":"6150_CR30","doi-asserted-by":"publisher","first-page":"37","DOI":"10.1016\/j.inffus.2022.11.022","volume":"92","author":"K Kim","year":"2023","unstructured":"Kim K, Park S (2023) Aobert: All-modalities-in-one Bert for multimodal sentiment analysis. Inf Fusion 92:37\u201345. https:\/\/doi.org\/10.1016\/j.inffus.2022.11.022","journal-title":"Inf Fusion"},{"key":"6150_CR31","doi-asserted-by":"publisher","first-page":"208","DOI":"10.1016\/j.ins.2023.01.116","volume":"628","author":"J Wang","year":"2023","unstructured":"Wang J, Wang S, Lin M, Xu Z, Guo W (2023) Learning speaker-independent multimodal representation for sentiment analysis. Inf Sci 628:208\u2013225. https:\/\/doi.org\/10.1016\/j.ins.2023.01.116","journal-title":"Inf Sci"},{"key":"6150_CR32","doi-asserted-by":"publisher","first-page":"109259","DOI":"10.1016\/j.patcog.2022.109259","volume":"136","author":"D Wang","year":"2023","unstructured":"Wang D, Guo X, Tian Y, Liu J, He L, Luo X (2023) Tetfn: A text enhanced transformer fusion network for multimodal sentiment analysis. Pattern Recognit 136:109259. https:\/\/doi.org\/10.1016\/j.patcog.2022.109259","journal-title":"Pattern Recognit"},{"key":"6150_CR33","doi-asserted-by":"publisher","first-page":"119125","DOI":"10.1016\/j.ins.2023.119125","volume":"641","author":"Z Tang","year":"2023","unstructured":"Tang Z, Xiao Q, Zhou X, Li Y, Chen C, Li K (2023) Learning discriminative multi-relation representations for multimodal sentiment analysis. Inf Sci 641:119125. https:\/\/doi.org\/10.1016\/j.ins.2023.119125","journal-title":"Inf Sci"},{"key":"6150_CR34","doi-asserted-by":"publisher","unstructured":"Issa B, Jasser MB, Chua HN, Hamzah M (2023) A comparative study on embedding models for keyword extraction using keybert method. In: ICSET 2023 - 2023 IEEE 13th international conference on system engineering and technology, Proceeding, pp 40\u201345. https:\/\/doi.org\/10.1109\/ICSET59111.2023.10295108","DOI":"10.1109\/ICSET59111.2023.10295108"},{"key":"6150_CR35","doi-asserted-by":"crossref","unstructured":"Loper E, Bird S (2002) Nltk: The natural language toolkit. In: ACL-02 Workshop on effective tools and methodologies for teaching natural language processing and computational linguistics, Proceedings, pp 63\u201370","DOI":"10.3115\/1118108.1118117"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-024-06150-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-024-06150-1\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-024-06150-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,1,30]],"date-time":"2025-01-30T16:02:00Z","timestamp":1738252920000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-024-06150-1"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,12,19]]},"references-count":35,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,2]]}},"alternative-id":["6150"],"URL":"https:\/\/doi.org\/10.1007\/s10489-024-06150-1","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,12,19]]},"assertion":[{"value":"4 December 2024","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"19 December 2024","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interests"}},{"value":"The authors confirm that there were no human research participants involved in this study and that the research received approval from the university\u2019s ethical committee.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical and informed consent for data used"}}],"article-number":"188"}}