{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,17]],"date-time":"2026-07-17T18:02:50Z","timestamp":1784311370536,"version":"3.55.0"},"reference-count":49,"publisher":"Elsevier BV","license":[{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-017"},{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"},{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-012"},{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2026,11,1]],"date-time":"2026-11-01T00:00:00Z","timestamp":1793491200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-004"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["Information Sciences"],"published-print":{"date-parts":[[2026,11]]},"DOI":"10.1016\/j.ins.2026.123792","type":"journal-article","created":{"date-parts":[[2026,6,15]],"date-time":"2026-06-15T16:30:49Z","timestamp":1781541049000},"page":"123792","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":0,"special_numbering":"C","title":["Pseudo adversarial alignment and preference decorrelation model for multimodal recommendation"],"prefix":"10.1016","volume":"755","author":[{"given":"Tao","family":"Xiong","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Mingming","family":"Yang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-6908-8018","authenticated-orcid":false,"given":"Wenming","family":"Cao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Man","family":"Li","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Wenda","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0935-5890","authenticated-orcid":false,"given":"Zhiwen","family":"Yu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-1530-7529","authenticated-orcid":false,"given":"Hau-San","family":"Wong","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/j.ins.2026.123792_bib0005","series-title":"Proceedings of the 22nd Annual International ACM SIGIR Conference on Research and Development in Information Retrieval","first-page":"230","article-title":"An algorithmic framework for performing collaborative filtering","author":"Herlocker","year":"1999"},{"key":"10.1016\/j.ins.2026.123792_bib0010","author":"Rendle"},{"key":"10.1016\/j.ins.2026.123792_bib0015","series-title":"Proceedings of the 43rd International ACM SIGIR Conference on Research and Development in Information Retrieval","first-page":"639","article-title":"Lightgcn: simplifying and powering graph convolution network for recommendation","author":"He","year":"2020"},{"key":"10.1016\/j.ins.2026.123792_bib0020","series-title":"Proceedings of the 39th International Conference on Data Engineering","first-page":"1247","article-title":"Layer-refined graph convolutional networks for recommendation","author":"Zhou","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0025","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","article-title":"VBPR: visual Bayesian personalized ranking from implicit feedback","volume":"vol. 30","author":"He","year":"2016"},{"key":"10.1016\/j.ins.2026.123792_bib0030","doi-asserted-by":"crossref","first-page":"2901","DOI":"10.1007\/s10489-020-01703-6","article-title":"Multi-view visual Bayesian personalized ranking for restaurant recommendation","volume":"50","author":"Zhang","year":"2020","journal-title":"Appl. Intell."},{"key":"10.1016\/j.ins.2026.123792_bib0035","series-title":"Proceedings of the 41st International ACM SIGIR Conference on Research and Development in Information Retrieval","first-page":"981","article-title":"Graphcar: content-aware multimedia recommendation with graph autoencoder","author":"Xu","year":"2018"},{"key":"10.1016\/j.ins.2026.123792_bib0040","doi-asserted-by":"crossref","first-page":"22","DOI":"10.1016\/j.ins.2021.09.006","article-title":"A two-stage embedding model for recommendation with multimodal auxiliary information","volume":"582","author":"Ni","year":"2022","journal-title":"Inf. Sci."},{"key":"10.1016\/j.ins.2026.123792_bib0045","series-title":"Proceedings of the 46th International ACM SIGIR Conference on Research and Development in Information Retrieval","first-page":"1508","article-title":"Lightgt: a light graph transformer for multimedia recommendation","author":"Wei","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0050","series-title":"Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval","first-page":"1807","article-title":"Multi-modal graph contrastive learning for micro-video recommendation","author":"Yi","year":"2022"},{"key":"10.1016\/j.ins.2026.123792_bib0055","doi-asserted-by":"crossref","DOI":"10.1016\/j.ins.2023.119760","article-title":"Self-supervised multimodal graph convolutional network for collaborative filtering","volume":"653","author":"Kim","year":"2024","journal-title":"Inf. Sci."},{"key":"10.1016\/j.ins.2026.123792_bib0060","series-title":"Proceedings of the 32nd ACM International Conference on Multimedia","first-page":"7591","article-title":"Diffmm: multi-modal diffusion model for recommendation","author":"Jiang","year":"2024"},{"key":"10.1016\/j.ins.2026.123792_bib0065","series-title":"Proceedings of the 33rd ACM International Conference on Information and Knowledge Management","first-page":"1503","article-title":"Alignrec: aligning and training in multimodal recommendations","author":"Liu","year":"2024"},{"key":"10.1016\/j.ins.2026.123792_bib0070","doi-asserted-by":"crossref","first-page":"855","DOI":"10.1109\/TKDE.2019.2893638","article-title":"Adversarial training towards robust multimedia recommender system","volume":"32","author":"Tang","year":"2019","journal-title":"IEEE Trans. Knowl. Data Eng."},{"key":"10.1016\/j.ins.2026.123792_bib0075","doi-asserted-by":"crossref","first-page":"8842","DOI":"10.1109\/TKDE.2024.3424268","article-title":"Multimodal graph causal embedding for multimedia-based recommendation","volume":"36","author":"Li","year":"2024","journal-title":"IEEE Trans. Knowl. Data Eng."},{"key":"10.1016\/j.ins.2026.123792_bib0080","doi-asserted-by":"crossref","DOI":"10.1016\/j.ins.2023.119815","article-title":"Graph neural networks with deep mutual learning for designing multi-modal recommendation systems","volume":"654","author":"Li","year":"2024","journal-title":"Inf. Sci."},{"key":"10.1016\/j.ins.2026.123792_bib0085","first-page":"1","article-title":"Multimodal recommender systems: a survey","volume":"57","author":"Liu","year":"2024","journal-title":"ACM Comput. Surv."},{"key":"10.1016\/j.ins.2026.123792_bib0090","series-title":"Proceedings of the 27th ACM International Conference on Multimedia","first-page":"1437","article-title":"MMGCN: multi-modal graph convolution network for personalized recommendation of micro-video","author":"Wei","year":"2019"},{"key":"10.1016\/j.ins.2026.123792_bib0095","series-title":"Proceedings of the 28th ACM International Conference on Multimedia","first-page":"3541","article-title":"Graph-refined convolutional network for multimedia recommendation with implicit feedback","author":"Wei","year":"2020"},{"key":"10.1016\/j.ins.2026.123792_bib0100","doi-asserted-by":"crossref","first-page":"1074","DOI":"10.1109\/TMM.2021.3138298","article-title":"Dualgnn: dual graph neural network for multimedia recommendation","volume":"25","author":"Wang","year":"2021","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.ins.2026.123792_bib0105","series-title":"Proceedings of the 30th ACM International Conference on Multimedia","first-page":"376","article-title":"Learning hybrid behavior patterns for multimedia recommendation","author":"Mu","year":"2022"},{"key":"10.1016\/j.ins.2026.123792_bib0110","series-title":"Proceedings of the 29th ACM International Conference on Multimedia","first-page":"3872","article-title":"Mining latent structures for multimedia recommendation","author":"Zhang","year":"2021"},{"key":"10.1016\/j.ins.2026.123792_bib0115","doi-asserted-by":"crossref","first-page":"9154","DOI":"10.1109\/TKDE.2022.3221949","article-title":"Latent structure mining with contrastive modality fusion for multimedia recommendation","volume":"35","author":"Zhang","year":"2022","journal-title":"IEEE Trans. Knowl. Data Eng."},{"key":"10.1016\/j.ins.2026.123792_bib0120","series-title":"Proceedings of the 31st ACM International Conference on Multimedia","first-page":"935","article-title":"A tale of two graphs: freezing and denoising graph structures for multimodal recommendation","author":"Zhou","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0125","series-title":"Proceedings of the 31st ACM International Conference on Multimedia","first-page":"6576","article-title":"Multi-view graph convolutional network for multimedia recommendation","author":"Yu","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0130","series-title":"ECAI 2023","first-page":"3123","article-title":"Enhancing dyadic relations with homogeneous graphs for multimodal recommendation","author":"Zhou","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0135","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"8454","article-title":"Lgmrec: local and global graph learning for multimodal recommendation","volume":"vol. 38","author":"Guo","year":"2024"},{"key":"10.1016\/j.ins.2026.123792_bib0140","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"11790","article-title":"Modality-independent graph neural networks with global transformers for multimodal recommendation","volume":"vol. 39","author":"Hu","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0145","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"13096","article-title":"Mind individual information! Principal graph learning for multimedia recommendation","volume":"vol. 39","author":"Yu","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0150","series-title":"Proceedings of the 48th International ACM SIGIR Conference on Research and Development in Information Retrieval","first-page":"1830","article-title":"COHESION: composite graph convolutional network with dual-stage fusion for multimodal recommendation","author":"Xu","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0155","doi-asserted-by":"crossref","first-page":"9052","DOI":"10.1109\/TPAMI.2024.3415112","article-title":"A survey on self-supervised learning: algorithms, applications, and future trends","volume":"46","author":"Gui","year":"2024","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.ins.2026.123792_bib0160","doi-asserted-by":"crossref","first-page":"5107","DOI":"10.1109\/TMM.2022.3187556","article-title":"Self-supervised learning for multimedia recommendation","volume":"25","author":"Tao","year":"2022","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.ins.2026.123792_bib0165","doi-asserted-by":"crossref","first-page":"9343","DOI":"10.1109\/TMM.2023.3251108","article-title":"Multimodal graph contrastive learning for multimedia-based recommendation","volume":"25","author":"Liu","year":"2023","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.ins.2026.123792_bib0170","series-title":"Proceedings of the ACM Web Conference","first-page":"790","article-title":"Multi-modal self-supervised learning for recommendation","author":"Wei","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0175","series-title":"Proceedings of the ACM Web Conference","first-page":"845","article-title":"Bootstrap latent representations for multi-modal recommendation","author":"Zhou","year":"2023"},{"key":"10.1016\/j.ins.2026.123792_bib0180","doi-asserted-by":"crossref","first-page":"8849","DOI":"10.1109\/TMM.2024.3382889","article-title":"SPACE: self-supervised dual preference enhancing network for multimodal recommendation","volume":"26","author":"Guo","year":"2024","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.ins.2026.123792_bib0185","series-title":"Proceedings of the 33rd ACM International Conference on Multimedia","first-page":"6316","article-title":"Refining contrastive learning and homography relations for multi-modal recommendation","author":"Ma","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0190","doi-asserted-by":"crossref","first-page":"18587","DOI":"10.1109\/TNNLS.2025.3583509","article-title":"Diffcl: a diffusion-based contrastive learning framework with semantic alignment for multimodal recommendations","volume":"36","author":"Song","year":"2025","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"key":"10.1016\/j.ins.2026.123792_bib0195","first-page":"1","article-title":"LPIC: learnable prompts and ID-guided contrastive learning for multimodal recommendation","volume":"21","author":"Liu","year":"2025","journal-title":"ACM Transactions on Multimedia Computing, Communications and Applications"},{"key":"10.1016\/j.ins.2026.123792_bib0200","series-title":"Proceedings of the 30th ACM SIGKDD Conference on Knowledge Discovery and Data Mining","first-page":"3645","article-title":"Improving multi-modal recommender systems by denoising and aligning multi-modal content and user feedback","author":"Xv","year":"2024"},{"key":"10.1016\/j.ins.2026.123792_bib0205","series-title":"Proceedings of the 34th ACM International Conference on Information and Knowledge Management","first-page":"2493","article-title":"Modality alignment with multi-scale bilateral attention for multimodal recommendation","author":"Ren","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0210","doi-asserted-by":"crossref","DOI":"10.1016\/j.inffus.2025.103472","article-title":"Dual-layer cross-modal alignment recommendation based on the diffusion model","volume":"125","author":"Xiu","year":"2026","journal-title":"Inf. Fusion"},{"key":"10.1016\/j.ins.2026.123792_bib0215","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"12908","article-title":"Mentor: multi-level self-supervised learning for multimodal recommendation","volume":"vol. 39","author":"Xu","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0220","doi-asserted-by":"crossref","DOI":"10.1016\/j.aei.2026.104462","article-title":"Dual-space feature representation learning network for multimodal recommender systems","volume":"72","author":"Dang","year":"2026","journal-title":"Adv. Eng. Inform."},{"key":"10.1016\/j.ins.2026.123792_bib0225","series-title":"Proceedings of the 25th International Conference on World Wide Web","first-page":"507","article-title":"Ups and downs: modeling the visual evolution of fashion trends with one-class collaborative filtering","author":"He","year":"2016"},{"key":"10.1016\/j.ins.2026.123792_bib0230","author":"Reimers"},{"key":"10.1016\/j.ins.2026.123792_bib0235","series-title":"Proceedings of the Eighteenth ACM International Conference on Web Search and Data Mining","first-page":"773","article-title":"Spectrum-based modality representation fusion graph convolutional network for multimodal recommendation","author":"Ong","year":"2025"},{"key":"10.1016\/j.ins.2026.123792_bib0240","doi-asserted-by":"crossref","DOI":"10.1016\/j.knosys.2025.113496","article-title":"Discrepancy learning guided hierarchical fusion network for multi-modal recommendation","volume":"317","author":"Dang","year":"2025","journal-title":"Knowl.-based Syst."},{"key":"10.1016\/j.ins.2026.123792_bib0245","series-title":"Proceedings of the 5th ACM International Conference on Multimedia in Asia Workshops","first-page":"1","article-title":"Mmrec: simplifying multimodal recommendation","author":"Zhou","year":"2023"}],"container-title":["Information Sciences"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0020025526007231?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0020025526007231?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2026,7,17]],"date-time":"2026-07-17T17:04:08Z","timestamp":1784307848000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0020025526007231"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,11]]},"references-count":49,"alternative-id":["S0020025526007231"],"URL":"https:\/\/doi.org\/10.1016\/j.ins.2026.123792","relation":{},"ISSN":["0020-0255"],"issn-type":[{"value":"0020-0255","type":"print"}],"subject":[],"published":{"date-parts":[[2026,11]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"Pseudo adversarial alignment and preference decorrelation model for multimodal recommendation","name":"articletitle","label":"Article Title"},{"value":"Information Sciences","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1016\/j.ins.2026.123792","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"\u00a9 2026 Elsevier Inc. All rights are reserved, including those for text and data mining, AI training, and similar technologies.","name":"copyright","label":"Copyright"}],"article-number":"123792"}}