{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,11]],"date-time":"2026-05-11T12:09:42Z","timestamp":1778501382574,"version":"3.51.4"},"publisher-location":"Singapore","reference-count":34,"publisher":"Springer Nature Singapore","isbn-type":[{"value":"9789819203680","type":"print"},{"value":"9789819203697","type":"electronic"}],"license":[{"start":{"date-parts":[[2026,1,1]],"date-time":"2026-01-01T00:00:00Z","timestamp":1767225600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,1,1]],"date-time":"2026-01-01T00:00:00Z","timestamp":1767225600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2026]]},"DOI":"10.1007\/978-981-92-0369-7_31","type":"book-chapter","created":{"date-parts":[[2026,5,11]],"date-time":"2026-05-11T11:23:35Z","timestamp":1778498615000},"page":"491-507","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Aspect-Oriented Prompt with\u00a0Adaptive Cross-Modal Fusion for\u00a0Multimodal Sentiment Analysis"],"prefix":"10.1007","author":[{"given":"Jianwei","family":"Zhang","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yutian","family":"Li","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yu","family":"Hu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Xudong","family":"Mao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Fuqiang","family":"Yu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Lap-Kei","family":"Lee","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Fu Lee","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zhenguo","family":"Yang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,5,12]]},"reference":[{"key":"31_CR1","doi-asserted-by":"crossref","unstructured":"Anderson, P., et al.: Bottom-up and top-down attention for image captioning and visual question answering. In: CVPR, pp. 6077\u20136086. Computer Vision Foundation\/IEEE Computer Society (2018)","DOI":"10.1109\/CVPR.2018.00636"},{"key":"31_CR2","doi-asserted-by":"crossref","unstructured":"Cao, Y., Bin, J., Hamari, J., Blasch, E., Liu, Z.: Multimodal object detection by channel switching and spatial attention. In: CVPR, pp. 403\u2013411 (2023)","DOI":"10.1109\/CVPRW59228.2023.00046"},{"key":"31_CR3","doi-asserted-by":"crossref","unstructured":"Chen, C.R., Fan, Q., Panda, R.: CrossVit: cross-attention multi-scale vision transformer for image classification. In: ICCV, pp. 347\u2013356. IEEE (2021)","DOI":"10.1109\/ICCV48922.2021.00041"},{"key":"31_CR4","doi-asserted-by":"crossref","unstructured":"Chen, L., Yan, X., Xiao, J., Zhang, H., Pu, S., Zhuang, Y.: Counterfactual samples synthesizing for robust visual question answering. In: CVPR, pp. 10797\u201310806. Computer Vision Foundation\/IEEE (2020)","DOI":"10.1109\/CVPR42600.2020.01081"},{"key":"31_CR5","doi-asserted-by":"crossref","unstructured":"Chen, Z., Zhu, Z., Xu, W., Zhang, Y., Wu, X., Zheng, Y.: Aspects are Anchors: towards multimodal aspect-based sentiment analysis via aspect-driven alignment and refinement. In: MM, pp. 2292\u20132300 (2024)","DOI":"10.1145\/3664647.3681189"},{"key":"31_CR6","doi-asserted-by":"crossref","unstructured":"Cheng, Z., et al.: Emotion-llama: Multimodal emotion recognition and reasoning with instruction tuning. In: Globersons, A., et al. (eds.) NeurIPS (2024)","DOI":"10.52202\/079017-3518"},{"key":"31_CR7","doi-asserted-by":"crossref","unstructured":"Devlin, J., Chang, M., Lee, K., Toutanova, K.: BERT: pre-training of deep bidirectional transformers for language understanding. In: NAACL, pp. 4171\u20134186 (2019)","DOI":"10.18653\/v1\/N19-1423"},{"key":"31_CR8","doi-asserted-by":"crossref","unstructured":"Gao, Z., Hu, D., Jiang, X., Lu, H., Shen, H.T., Xu, X.: Enhanced experts with uncertainty-aware routing for multimodal sentiment analysis. In: MM, pp. 9650\u20139659 (2024)","DOI":"10.1145\/3664647.3680949"},{"key":"31_CR9","doi-asserted-by":"crossref","unstructured":"Han, Z., Hu, M., Bai, Y., Wang, X., Luo, B.: DEQA: descriptions enhanced question-answering framework for multimodal aspect-based sentiment analysis. In: Walsh, T., Shah, J., Kolter, Z. (eds.) AAAI, pp. 23987\u201323995. AAAI Press (2025)","DOI":"10.1609\/aaai.v39i22.34572"},{"key":"31_CR10","doi-asserted-by":"crossref","unstructured":"Huang, J., Ji, Y., Qin, Z., Yang, Y., Shen, H.T.: Dominant single-modal supplementary fusion (SIMSUF) for multimodal sentiment analysis. IEEE Trans. Multim. 26, 8383\u20138394 (2024)","DOI":"10.1109\/TMM.2023.3344358"},{"key":"31_CR11","doi-asserted-by":"crossref","unstructured":"Jiang, X., Tang, H., Gao, J., Du, X., He, S., Li, Z.: Delving into multimodal prompting for fine-grained visual classification. In: AAAI, pp. 2570\u20132578 (2024)","DOI":"10.1609\/aaai.v38i3.28034"},{"key":"31_CR12","doi-asserted-by":"crossref","unstructured":"Ju, X., et al.: Joint multi-modal aspect-sentiment analysis with auxiliary cross-modal relation detection. In: Moens, M., Huang, X., Specia, L., Yih, S.W. (eds.) EMNLP, pp. 4395\u20134405. Association for Computational Linguistics (2021)","DOI":"10.18653\/v1\/2021.emnlp-main.360"},{"key":"31_CR13","doi-asserted-by":"crossref","unstructured":"Khan, Z., Fu, Y.: Exploiting BERT for multimodal target sentiment classification through input space translation. In: Shen, H.T., Zhuang, Y., Smith, J.R., Yang, Y., C\u00e9sar, P., Metze, F., Prabhakaran, B. (eds.) MM, pp. 3034\u20133042. ACM (2021)","DOI":"10.1145\/3474085.3475692"},{"key":"31_CR14","doi-asserted-by":"crossref","unstructured":"Lewis, M., et al.: BART: denoising sequence-to-sequence pre-training for natural language generation, translation, and comprehension. In: ACL, pp. 7871\u20137880 (2020)","DOI":"10.18653\/v1\/2020.acl-main.703"},{"key":"31_CR15","doi-asserted-by":"crossref","unstructured":"Ling, Y., Yu, J., Xia, R.: Vision-language pre-training for multimodal aspect-based sentiment analysis. In: Muresan, S., Nakov, P., Villavicencio, A. (eds.) ACL, pp. 2149\u20132159. Association for Computational Linguistics (2022)","DOI":"10.18653\/v1\/2022.acl-long.152"},{"key":"31_CR16","doi-asserted-by":"crossref","unstructured":"Liu, H., Li, C., Li, Y., Lee, Y.J.: Improved baselines with visual instruction tuning. In: CVPR, pp. 26286\u201326296 (2024)","DOI":"10.1109\/CVPR52733.2024.02484"},{"key":"31_CR17","doi-asserted-by":"crossref","unstructured":"Mitra, C., Huang, B., Darrell, T., Herzig, R.: Compositional chain-of-thought prompting for large multimodal models. In: CVPR, pp. 14420\u201314431 (2024)","DOI":"10.1109\/CVPR52733.2024.01367"},{"key":"31_CR18","doi-asserted-by":"crossref","unstructured":"Peng, T., Li, Z., Wang, P., Zhang, L., Zhao, H.: A novel energy based model mechanism for multi-modal aspect-based sentiment analysis. In: AAAI, pp. 18869\u201318878 (2024)","DOI":"10.1609\/aaai.v38i17.29852"},{"key":"31_CR19","unstructured":"Su, W., et al.: VL-BERT: pre-training of generic visual-linguistic representations. In: ICLR (2020). OpenReview.net"},{"key":"31_CR20","doi-asserted-by":"crossref","unstructured":"Sun, L., Wang, J., Zhang, K., Su, Y., Weng, F.: RPBert: a text-image relation propagation-based BERT model for multimodal NER. In: AAAI, pp. 13860\u201313868 (2021)","DOI":"10.1609\/aaai.v35i15.17633"},{"key":"31_CR21","doi-asserted-by":"crossref","unstructured":"Vinyals, O., Toshev, A., Bengio, S., Erhan, D.: Show and tell: a neural image caption generator. In: CVPR (2015)","DOI":"10.1109\/CVPR.2015.7298935"},{"key":"31_CR22","doi-asserted-by":"publisher","unstructured":"Wu, H., Cheng, S., Wang, J., Li, S., Chi, L.: Multimodal aspect extraction with region-aware alignment network. In: NLPCC. LNCS, vol. 12430, pp. 145\u2013156 (2020). https:\/\/doi.org\/10.1007\/978-3-030-60450-9_12","DOI":"10.1007\/978-3-030-60450-9_12"},{"key":"31_CR23","doi-asserted-by":"crossref","unstructured":"Wu, Z., Zheng, C., Cai, Y., Chen, J., Leung, H., Li, Q.: Multimodal representation with embedded visual guiding objects for named entity recognition in social media posts. In: MM, pp. 1038\u20131046 (2020)","DOI":"10.1145\/3394171.3413650"},{"key":"31_CR24","doi-asserted-by":"crossref","unstructured":"Xiao, L., Wu, X., Xu, J., Li, W., Jin, C., He, L.: Atlantis: aesthetic-oriented multiple granularities fusion network for joint multimodal aspect-based sentiment analysis. Inf. Fusion 106, 102304 (2024)","DOI":"10.1016\/j.inffus.2024.102304"},{"key":"31_CR25","unstructured":"Xu, K., et al.: Show, attend and tell: neural image caption generation with visual attention. In: Bach, F.R., Blei, D.M. (eds.) ICML. JMLR Workshop and Conference Proceedings, vol.\u00a037, pp. 2048\u20132057. JMLR.org (2015)"},{"key":"31_CR26","doi-asserted-by":"crossref","unstructured":"Yu, J., Jiang, J.: Adapting BERT for target-oriented multimodal sentiment classification. In: Kraus, S. (ed.) IJCAI, pp. 5408\u20135414 (2019). ijcai.org","DOI":"10.24963\/ijcai.2019\/751"},{"key":"31_CR27","doi-asserted-by":"crossref","unstructured":"Yu, J., Jiang, J., Xia, R.: Entity-sensitive attention and fusion network for entity-level multimodal sentiment classification. IEEE ACM Trans. Audio Speech Lang. Process. 28, 429\u2013439 (2020)","DOI":"10.1109\/TASLP.2019.2957872"},{"key":"31_CR28","doi-asserted-by":"crossref","unstructured":"Yu, J., Jiang, J., Yang, L., Xia, R.: Improving multimodal named entity recognition via entity span detection with unified multimodal transformer. In: ACL, pp. 3342\u20133352 (2020)","DOI":"10.18653\/v1\/2020.acl-main.306"},{"key":"31_CR29","doi-asserted-by":"crossref","unstructured":"Yu, Z., Yu, J., Fan, J., Tao, D.: Multi-modal factorized bilinear pooling with co-attention learning for visual question answering. In: ICCV, pp. 1839\u20131848. IEEE Computer Society (2017)","DOI":"10.1109\/ICCV.2017.202"},{"key":"31_CR30","doi-asserted-by":"crossref","unstructured":"Zhang, D., Wei, S., Li, S., Wu, H., Zhu, Q., Zhou, G.: Multi-modal graph fusion for named entity recognition with targeted visual guidance. In: AAAI, pp. 14347\u201314355 (2021)","DOI":"10.1609\/aaai.v35i16.17687"},{"key":"31_CR31","doi-asserted-by":"crossref","unstructured":"Zhang, Y., Dong, Y., Zhang, S., Min, T., Su, H., Zhu, J.: Exploring the transferability of visual prompting for multimodal large language models. In: CVPR, pp. 26552\u201326562 (2024)","DOI":"10.1109\/CVPR52733.2024.02508"},{"key":"31_CR32","doi-asserted-by":"crossref","unstructured":"Zhou, Q., et al.: Token-level contrastive learning with modality-aware prompting for multimodal intent recognition. In: AAAI, pp. 17114\u201317122 (2024)","DOI":"10.1609\/aaai.v38i15.29656"},{"key":"31_CR33","doi-asserted-by":"crossref","unstructured":"Zhou, R., Guo, W., Liu, X., Yu, S., Zhang, Y., Yuan, X.: AOM: detecting aspect-oriented information for multimodal aspect-based sentiment analysis. In: Rogers, A., Boyd-Graber, J.L., Okazaki, N. (eds.) ACL, pp. 8184\u20138196 (2023)","DOI":"10.18653\/v1\/2023.findings-acl.519"},{"key":"31_CR34","unstructured":"Zhu, L., Sun, H., Gao, Q., Yi, T., He, L.: Joint multimodal aspect sentiment analysis with aspect enhancement and syntactic adaptive learning. In: IJCAI, pp. 6678\u20136686. ijcai.org (2024)"}],"container-title":["Lecture Notes in Computer Science","Database Systems for Advanced Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-981-92-0369-7_31","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,5,11]],"date-time":"2026-05-11T11:24:09Z","timestamp":1778498649000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-981-92-0369-7_31"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026]]},"ISBN":["9789819203680","9789819203697"],"references-count":34,"URL":"https:\/\/doi.org\/10.1007\/978-981-92-0369-7_31","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026]]},"assertion":[{"value":"12 May 2026","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Disclosure of Interests"}},{"value":"DASFAA","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"International Conference on Database Systems for Advanced Applications","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Jeju","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Korea (Republic of)","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2026","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"27 April 2026","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"30 April 2026","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"31","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"dasfaa2026","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/dasfaa2026.github.io\/index.html","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}