{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,8]],"date-time":"2026-07-08T17:17:48Z","timestamp":1783531068217,"version":"3.55.0"},"reference-count":59,"publisher":"Elsevier BV","license":[{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-017"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-012"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-004"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100006385","name":"Chengdu University of Technology","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100006385","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["Expert Systems with Applications"],"published-print":{"date-parts":[[2026,12]]},"DOI":"10.1016\/j.eswa.2026.133254","type":"journal-article","created":{"date-parts":[[2026,6,13]],"date-time":"2026-06-13T01:03:14Z","timestamp":1781312594000},"page":"133254","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":0,"special_numbering":"PB","title":["Breaking sequential bias: Importance-guided temporal multimodal recommendation via state space models"],"prefix":"10.1016","volume":"331","author":[{"ORCID":"https:\/\/orcid.org\/0009-0000-9092-0668","authenticated-orcid":false,"given":"Kang","family":"Zhang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-7194-9318","authenticated-orcid":false,"given":"Quan","family":"Wen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yujian","family":"Huang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-2295-1328","authenticated-orcid":false,"given":"Yanmei","family":"Hu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Ruixing","family":"Huang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Na","family":"Dong","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xiaomeng","family":"Yang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Shuyi","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/j.eswa.2026.133254_bib0001","unstructured":"Betello, F., Siciliano, F., Mishra, P., & Silvestri, F. (2023). Investigating the robustness of sequential recommender systems against training data perturbations: An empirical study(vol. abs\/2307.13165). arXiv: 2307.13165. https:\/\/api.semanticscholar.org\/CorpusID:260154990."},{"key":"10.1016\/j.eswa.2026.133254_sbref0002","series-title":"Proceedings of the 32nd acm international conference on information and knowledge management","article-title":"Multi-modal mixture of experts representation learning for sequential recommendation","author":"Bian","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_sbref0003","doi-asserted-by":"crossref","DOI":"10.1007\/s11063-024-11438-x","article-title":"A robust sequential recommendation model based on multiple feedback behavior denoising and trusted neighbors","volume":"56","author":"Cai","year":"2024","journal-title":"Neural Processing Letters"},{"key":"10.1016\/j.eswa.2026.133254_sbref0004","series-title":"Proceedings of the 16th acm conference on recommender systems","article-title":"Denoising self-attentive sequential recommendation","author":"Chen","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_sbref0005","series-title":"Knowledge science, engineering and management","article-title":"Mmea: Entity alignment for multi-modal knowledge graph","author":"Chen","year":"2020"},{"key":"10.1016\/j.eswa.2026.133254_sbref0006","series-title":"Proceedings of the 28th acm sigkdd conference on knowledge discovery and data mining","article-title":"Multi-modal siamese network for entity alignment","author":"Chen","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_sbref0007","series-title":"Advances in neural information processing systems 37","article-title":"Tackling uncertain correspondences for multi-modal entity alignment","author":"Chen","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_bib0008","unstructured":"Dao, T., & Gu, A. (2024). Transformers are ssms: Generalized models and efficient algorithms through structured state space duality. arxiv: 2405.21060. https:\/\/api.semanticscholar.org\/CorpusID:270199762."},{"key":"10.1016\/j.eswa.2026.133254_sbref0009","series-title":"North american chapter of the association for computational linguistics","article-title":"Bert: Pre-training of deep bidirectional transformers for language understanding","author":"Devlin","year":"2019"},{"key":"10.1016\/j.eswa.2026.133254_bib0010","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., Uszkoreit, J., & Houlsby, N. (2020). An image is worth 16x16 words: Transformers for image recognition at scale. ArXiv:2010.11929, https:\/\/api.semanticscholar.org\/CorpusID:225039882."},{"key":"10.1016\/j.eswa.2026.133254_sbref0011","series-title":"Proceedings of the 46th international acm sigir conference on research and development in information retrieval","article-title":"Frequency enhanced hybrid attention network for sequential recommendation","author":"Du","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_bib0012","doi-asserted-by":"crossref","unstructured":"Fan, H., Zhu, M., Hu, Y., Feng, H., He, Z., Liu, H., & Liu, Q. (2024). TiM4Rec: An efficient sequential recommendation model based on time-aware structured state space duality model. arxiv: 2409.16182https:\/\/api.semanticscholar.org\/CorpusID:272832506.","DOI":"10.1016\/j.neucom.2025.131270"},{"key":"10.1016\/j.eswa.2026.133254_bib0013","unstructured":"Gu, A., & Dao, T. (2023). Mamba: Linear-time sequence modeling with selective state spaces. https:\/\/api.semanticscholar.org\/CorpusID:265551773."},{"key":"10.1016\/j.eswa.2026.133254_bib0014","unstructured":"Gu, A., Goel, K., & R\u2019e, C. (2021). Efficiently modeling long sequences with structured state spaces. arxiv: 2111.00396https:\/\/api.semanticscholar.org\/CorpusID:240354066."},{"key":"10.1016\/j.eswa.2026.133254_sbref0015","series-title":"2025 Ieee international conference on acoustics, speech and signal processing (icassp)","first-page":"1","article-title":"M3rec: Selective state space models with mixture-of-modality experts for multi-modal sequential recommendation","author":"Guo","year":"2025"},{"key":"10.1016\/j.eswa.2026.133254_sbref0016","series-title":"Aaai conference on artificial intelligence","article-title":"VBPR: Visual bayesian personalized ranking from implicit feedback","author":"He","year":"2015"},{"key":"10.1016\/j.eswa.2026.133254_sbref0017","series-title":"Proceedings of the 43rd international acm sigir conference on research and development in information retrieval","article-title":"Lightgcn: Simplifying and powering graph convolution network for recommendation","author":"He","year":"2020"},{"key":"10.1016\/j.eswa.2026.133254_sbref0018","series-title":"Proceedings of the 26th international conference on world wide web","article-title":"Neural collaborative filtering","author":"He","year":"2017"},{"key":"10.1016\/j.eswa.2026.133254_sbref0019","article-title":"Session-based recommendations with recurrent neural networks","volume":"abs\/1511.06939","author":"Hidasi","year":"2015","journal-title":"CoRR"},{"key":"10.1016\/j.eswa.2026.133254_sbref0020","series-title":"Proceedings of the 28th acm sigkdd conference on knowledge discovery and data mining","article-title":"Towards universal sequence representation learning for recommender systems","author":"Hou","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_sbref0021","series-title":"Proceedings of the 32nd acm international conference on information and knowledge management","article-title":"Adaptive multi-modalities fusion in sequential recommendation systems","author":"Hu","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_sbref0022","series-title":"Proceedings of the 31st acm international conference on multimedia","article-title":"Online distillation-enhanced multi-modal transformer for sequential recommendation","author":"Ji","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_sbref0023","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1145\/3670995","article-title":"TriMLP: A foundational MLP-like architecture for sequential recommendation","volume":"42","author":"Jiang","year":"2024","journal-title":"ACM Transactions on Information Systems"},{"key":"10.1016\/j.eswa.2026.133254_sbref0024","first-page":"1","article-title":"Contrastive self-supervised learning in recommender systems: A survey","volume":"42","author":"Jing","year":"2023","journal-title":"ACM Transactions on Information Systems"},{"key":"10.1016\/j.eswa.2026.133254_sbref0025","series-title":"2018 Ieee international conference on data mining (icdm)","first-page":"197","article-title":"Self-attentive sequential recommendation","author":"Kang","year":"2018"},{"key":"10.1016\/j.eswa.2026.133254_sbref0026","series-title":"International conference on machine learning","article-title":"Dstagnn: Dynamic spatial-temporal aware graph neural network for traffic flow forecasting","author":"Lan","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_sbref0027","series-title":"Proceedings of the 29th acm sigkdd conference on knowledge discovery and data mining","article-title":"Text is all you need: Learning language representations for sequential recommendation","author":"Li","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_sbref0028","series-title":"Proceedings of the 13th international conference on web search and data mining","article-title":"Time interval aware self-attention for sequential recommendation","author":"Li","year":"2020"},{"key":"10.1016\/j.eswa.2026.133254_sbref0029","series-title":"Proceedings of the 17th acm international conference on web search and data mining","article-title":"Multi-sequence attentive user representation learning for side-information integrated sequential recommendation","author":"Lin","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_sbref0030","series-title":"Aaai conference on artificial intelligence","article-title":"Non-invasive self-attention for side information fusion in sequential recommendation","author":"Liu","year":"2021"},{"key":"10.1016\/j.eswa.2026.133254_bib0031","unstructured":"Liu, C., Lin, J., Wang, J., Liu, H., & Caverlee, J. (2024a). Mamba4Rec: Towards efficient sequential recommendation with selective state space models. https:\/\/api.semanticscholar.org\/CorpusID:268253535."},{"key":"10.1016\/j.eswa.2026.133254_bib0032","unstructured":"Liu, Z., Chen, Y., Li, J., Yu, P. S., McAuley, J., & Xiong, C. (2021b). Contrastive self-supervised sequential recommendation with robust augmentation. arxiv: 2108.06479. https:\/\/api.semanticscholar.org\/CorpusID:237091262."},{"key":"10.1016\/j.eswa.2026.133254_sbref0033","series-title":"Aaai conference on artificial intelligence","doi-asserted-by":"crossref","DOI":"10.5772\/intechopen.111293","article-title":"SIGMA: Selective gated mamba for sequential recommendation","author":"Liu","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_sbref0034","series-title":"Proceedings of the fifteenth acm international conference on web search and data mining","article-title":"Contrastive learning for representation degeneration problem in sequential recommendation","author":"Qiu","year":"2021"},{"key":"10.1016\/j.eswa.2026.133254_sbref0035","series-title":"Aaai conference on artificial intelligence","article-title":"An attentive inductive bias for sequential recommendation beyond the self-attention","author":"Shin","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_sbref0036","series-title":"Proceedings of the 28th acm international conference on information and knowledge management","article-title":"Bert4rec: Sequential recommendation with bidirectional encoder representations from transformer","author":"Sun","year":"2019"},{"key":"10.1016\/j.eswa.2026.133254_sbref0037","series-title":"Proceedings of the eleventh acm international conference on web search and data mining","article-title":"Personalized top-n sequential recommendation via convolutional sequence embedding","author":"Tang","year":"2018"},{"key":"10.1016\/j.eswa.2026.133254_bib0038","unstructured":"Tay, Y. T., Dehghani, M. D., Abnar, S. A., Shen, Y. S., Bahri, D. B., Pham, P. P., Rao, J. R., Yang, L. Y., Ruder, S. R., & Metzler, D. M. (2020). Long range arena: A benchmark for efficient transformers. arxiv: 2011.04006. https:\/\/api.semanticscholar.org\/CorpusID:260440449."},{"key":"10.1016\/j.eswa.2026.133254_sbref0039","series-title":"Companion proceedings of the acm web conference 2024","article-title":"Aligned side information fusion method for sequential recommendation","author":"Wang","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_sbref0040","first-page":"4555","article-title":"A survey on curriculum learning","volume":"44","author":"Wang","year":"2021","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"10.1016\/j.eswa.2026.133254_sbref0041","series-title":"Proceedings of the acm web conference 2023","article-title":"Multi-modal self-supervised learning for recommendation","author":"Wei","year":"2023"},{"key":"10.1016\/j.eswa.2026.133254_sbref0042","series-title":"Proceedings of the 27th acm international conference on multimedia","article-title":"MMGCN: Multi-modal graph convolution network for personalized recommendation of micro-video","author":"wei Wei","year":"2019"},{"key":"10.1016\/j.eswa.2026.133254_sbref0043","first-page":"1","article-title":"Learning robust sequential recommenders through confident soft labels","volume":"43","author":"Wu","year":"2023","journal-title":"ACM Transactions on Information Systems"},{"key":"10.1016\/j.eswa.2026.133254_sbref0044","series-title":"Proceedings of the 34th acm international conference on information and knowledge management","article-title":"Empowering denoising sequential recommendation with large language model embeddings","author":"Wu","year":"2025"},{"key":"10.1016\/j.eswa.2026.133254_sbref0045","series-title":"Proceedings of the 47th international acm sigir conference on research and development in information retrieval","article-title":"Afdgcf: Adaptive feature de-correlation graph collaborative filtering for recommendations","author":"Wu","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_bib0046","unstructured":"Xie, X., Sun, F., Liu, Z., Gao, J., Ding, B., & Cui, B. (2020). Contrastive pre-training for sequential recommendation. arxiv: 2010.14395. https:\/\/api.semanticscholar.org\/CorpusID:225076163."},{"key":"10.1016\/j.eswa.2026.133254_sbref0047","series-title":"Proceedings of the 45th international acm sigir conference on research and development in information retrieval","article-title":"Decoupled side information fusion for sequential recommendation","author":"Xie","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_sbref0048","series-title":"Aaai conference on artificial intelligence","article-title":"Wavelet enhanced adaptive frequency filter for sequential recommendation","author":"Xu","year":"2025"},{"key":"10.1016\/j.eswa.2026.133254_sbref0049","series-title":"Proceedings of the 33rd acm international conference on information and knowledge management","article-title":"Sequence-level semantic representation fusion for recommender systems","author":"Xu","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_sbref0050","series-title":"Aaai conference on artificial intelligence","article-title":"Harnessing multimodal large language models for multimodal sequential recommendation","author":"Ye","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_sbref0051","series-title":"Proceedings of the 31st acm international conference on information & knowledge management","article-title":"Hierarchical item inconsistency signal learning for sequence denoising in sequential recommendation","author":"Zhang","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_sbref0052","series-title":"Proceedings of the acm on web conference 2025","article-title":"Hierarchical time-aware mixture of experts for multi-modal sequential recommendation","author":"Zhang","year":"2025"},{"key":"10.1016\/j.eswa.2026.133254_sbref0053","series-title":"Proceedings of the 44th international acm sigir conference on research and development in information retrieval","article-title":"Causerec: Counterfactual user sequence synthesis for sequential recommendation","author":"Zhang","year":"2021"},{"key":"10.1016\/j.eswa.2026.133254_sbref0054","series-title":"International joint conference on artificial intelligence","article-title":"Feature-level deeper self-attention network for sequential recommendation","author":"Zhang","year":"2019"},{"key":"10.1016\/j.eswa.2026.133254_sbref0055","series-title":"SSD4Rec: A structured state space duality model for efficient sequential recommendation","author":"Zhang","year":"2024"},{"key":"10.1016\/j.eswa.2026.133254_sbref0056","series-title":"Proceedings of the 30th acm international conference on information & knowledge management","article-title":"RecBole: Towards a unified, comprehensive and efficient framework for recommendation algorithms","author":"Zhao","year":"2020"},{"key":"10.1016\/j.eswa.2026.133254_sbref0057","series-title":"Proceedings of the 29th acm international conference on information & knowledge management","article-title":"S3-rec: Self-supervised learning for sequential recommendation with mutual information maximization","author":"Zhou","year":"2020"},{"key":"10.1016\/j.eswa.2026.133254_sbref0058","series-title":"Proceedings of the acm web conference 2022","article-title":"Filter-enhanced mlp is all you need for sequential recommendation","author":"Zhou","year":"2022"},{"key":"10.1016\/j.eswa.2026.133254_bib0059","doi-asserted-by":"crossref","unstructured":"Zong, Z., Ma, B., Shen, D., Song, G., Shao, H., Jiang, D., Li, H., & Liu, Y. (2024). MoVA: Adapting mixture of vision experts to multimodal context. https:\/\/api.semanticscholar.org\/CorpusID:269282964.","DOI":"10.52202\/079017-3282"}],"container-title":["Expert Systems with Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0957417426021639?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0957417426021639?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2026,7,8]],"date-time":"2026-07-08T16:41:42Z","timestamp":1783528902000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0957417426021639"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,12]]},"references-count":59,"alternative-id":["S0957417426021639"],"URL":"https:\/\/doi.org\/10.1016\/j.eswa.2026.133254","relation":{},"ISSN":["0957-4174"],"issn-type":[{"value":"0957-4174","type":"print"}],"subject":[],"published":{"date-parts":[[2026,12]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"Breaking sequential bias: Importance-guided temporal multimodal recommendation via state space models","name":"articletitle","label":"Article Title"},{"value":"Expert Systems with Applications","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1016\/j.eswa.2026.133254","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"\u00a9 2026 Elsevier Ltd. All rights are reserved, including those for text and data mining, AI training, and similar technologies.","name":"copyright","label":"Copyright"}],"article-number":"133254"}}