{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,26]],"date-time":"2026-06-26T14:52:44Z","timestamp":1782485564102,"version":"3.54.5"},"reference-count":217,"publisher":"Elsevier BV","license":[{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-017"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-012"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2026,12,1]],"date-time":"2026-12-01T00:00:00Z","timestamp":1796083200000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-004"}],"funder":[{"DOI":"10.13039\/100009133","name":"Karlsruhe Institute of Technology","doi-asserted-by":"publisher","id":[{"id":"10.13039\/100009133","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["Information Fusion"],"published-print":{"date-parts":[[2026,12]]},"DOI":"10.1016\/j.inffus.2026.104543","type":"journal-article","created":{"date-parts":[[2026,6,14]],"date-time":"2026-06-14T17:35:52Z","timestamp":1781458552000},"page":"104543","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":0,"special_numbering":"C","title":["Toward a universal perception layer: A survey on sensor-agnostic advanced driver assistance systems"],"prefix":"10.1016","volume":"136","author":[{"ORCID":"https:\/\/orcid.org\/0009-0006-3391-277X","authenticated-orcid":false,"given":"Tim Alexander","family":"Bader","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0001-7079-9220","authenticated-orcid":false,"given":"Tim Dieter","family":"Eberhardt","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0003-4228-3205","authenticated-orcid":false,"given":"Tin Stribor","family":"Sohn","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-0579-4615","authenticated-orcid":false,"given":"Wilhelm","family":"Stork","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/j.inffus.2026.104543_bib0001","series-title":"The Case for an End-to-End Automotive-Software Platform","volume":"16","author":"Fletcher","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0002","article-title":"Deep learning in computer vision: a critical review of emerging techniques and application scenarios","volume":"6","author":"Chai","year":"2021","journal-title":"Mach. Learn. Appl."},{"key":"10.1016\/j.inffus.2026.104543_bib0003","series-title":"Future of Software Engineering (FOSE \u201907)","first-page":"55","article-title":"Software engineering for automotive systems: a roadmap","author":"Pretschner","year":"2007"},{"issue":"2","key":"10.1016\/j.inffus.2026.104543_bib0004","doi-asserted-by":"crossref","first-page":"356","DOI":"10.1109\/JPROC.2006.888386","article-title":"Engineering automotive software","volume":"95","author":"Broy","year":"2007","journal-title":"Proc. IEEE"},{"key":"10.1016\/j.inffus.2026.104543_bib0005","doi-asserted-by":"crossref","unstructured":"H. Reichert, L. Lang, K. R\u00f6sch, D. Bogdoll, K. Doll, B. Sick, H.-C. Reuss, C. Stiller, J.M. Z\u00f6llner, Towards sensor data abstraction of autonomous vehicle perception systems, (2021). 10.48550\/arXiv.2105.06896.","DOI":"10.1109\/ISC253183.2021.9562912"},{"key":"10.1016\/j.inffus.2026.104543_bib0006","series-title":"2025 IEEE Intelligent Vehicles Symposium (IV)","first-page":"1400","article-title":"Neural rendering for sensor adaptation in 3D object detection","author":"Embacher","year":"2025"},{"issue":"4","key":"10.1016\/j.inffus.2026.104543_bib0007","first-page":"94","article-title":"A review of sensor technologies for perception in automated driving","volume":"11","author":"Marti","year":"2019","journal-title":"IEEE Intell. Transp. Syst. Mag."},{"issue":"6","key":"10.1016\/j.inffus.2026.104543_bib0008","doi-asserted-by":"crossref","first-page":"2140","DOI":"10.3390\/s21062140","article-title":"Sensor and sensor fusion technology in autonomous vehicles: a review","volume":"21","author":"De Jong","year":"2021","journal-title":"Sensors"},{"key":"10.1016\/j.inffus.2026.104543_bib0009","doi-asserted-by":"crossref","DOI":"10.1016\/j.simpa.2022.100393","article-title":"OpenCalib: a multi-sensor calibration toolbox for autonomous driving","volume":"14","author":"Yan","year":"2022","journal-title":"Softw. Impacts"},{"key":"10.1016\/j.inffus.2026.104543_bib0010","unstructured":"On-Road Automated Driving (ORAD) Committee, Taxonomy and definitions for terms related to driving automation systems for on-road motor vehicles, (2021). 10.4271\/J3016_202104."},{"issue":"12","key":"10.1016\/j.inffus.2026.104543_bib0011","doi-asserted-by":"crossref","first-page":"6999","DOI":"10.1109\/TNNLS.2021.3084827","article-title":"A survey of convolutional neural networks: analysis, applications, and prospects","volume":"33","author":"Li","year":"2022","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0012","unstructured":"R. Lagani\u00e8re, Strategies and methods for sensor fusion, 2022. https:\/\/www.site.uottawa.ca\/research\/viva\/projects\/raddet\/index.html."},{"key":"10.1016\/j.inffus.2026.104543_bib0013","series-title":"Computer Vision","first-page":"983","article-title":"Pinhole camera model","author":"Sturm","year":"2021"},{"issue":"14","key":"10.1016\/j.inffus.2026.104543_bib0014","doi-asserted-by":"crossref","first-page":"14165","DOI":"10.1109\/JSEN.2022.3169805","article-title":"Discussion of novel filters and models for color space conversion","volume":"22","author":"Lelowicz","year":"2022","journal-title":"IEEE Sens. J."},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0015","doi-asserted-by":"crossref","first-page":"288","DOI":"10.2352\/issn.2169-2629.2020.28.46","article-title":"Optimization of automotive color filter arrays for traffic light color separation","volume":"28","author":"Weikl","year":"2020","journal-title":"Color Imaging Conf."},{"key":"10.1016\/j.inffus.2026.104543_bib0016","doi-asserted-by":"crossref","first-page":"1","DOI":"10.2352\/EI.2022.34.16.AVM-215","article-title":"Non-RGB color filter options and traffic signal detection capabilities","volume":"34","author":"Funatsu","year":"2022","journal-title":"Electron. Imaging"},{"issue":"4","key":"10.1016\/j.inffus.2026.104543_bib0017","doi-asserted-by":"crossref","first-page":"3638","DOI":"10.1109\/TITS.2023.3235057","article-title":"Surround-view fisheye camera perception for automated driving: overview, survey & challenges","volume":"24","author":"Kumar","year":"2023","journal-title":"IEEE Trans. Intell. Transp. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0018","doi-asserted-by":"crossref","first-page":"88","DOI":"10.1016\/j.imavis.2017.07.002","article-title":"Computer vision in automated parking systems: design, implementation and challenges","volume":"68","author":"Heimberger","year":"2017","journal-title":"Image Vis. Comput."},{"key":"10.1016\/j.inffus.2026.104543_bib0019","doi-asserted-by":"crossref","unstructured":"P. Li, X. Chen, S. Shen, Stereo R-CNN based 3D object detection for autonomous driving, 2019, pp. 7644\u20137652. https:\/\/openaccess.thecvf.com\/content_CVPR_2019\/html\/Li_Stereo_R-CNN_Based_3D_Object_Detection_for_Autonomous_Driving_CVPR_2019_paper.html.","DOI":"10.1109\/CVPR.2019.00783"},{"key":"10.1016\/j.inffus.2026.104543_bib0020","series-title":"2019 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"899","article-title":"Drivingstereo: a large-scale dataset for stereo matching in autonomous driving scenarios","author":"Yang","year":"2019"},{"key":"10.1016\/j.inffus.2026.104543_bib0021","series-title":"Modern Radar for Automotive Applications","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0022","series-title":"Unlocking radar sensor integration in vehicle: accelerated product development harnessing electromagnetic simulation","author":"Rao","year":"2024"},{"issue":"2","key":"10.1016\/j.inffus.2026.104543_bib0023","doi-asserted-by":"crossref","first-page":"36","DOI":"10.1109\/MSP.2016.2637700","article-title":"Advances in automotive radar: a framework on computationally efficient high-resolution frequency estimation","volume":"34","author":"Engels","year":"2017","journal-title":"IEEE Signal Process. Mag."},{"issue":"10","key":"10.1016\/j.inffus.2026.104543_bib0024","doi-asserted-by":"crossref","first-page":"135","DOI":"10.1109\/MCOM.2017.1700030","article-title":"Lidar system architectures and circuits","volume":"55","author":"Behroozpour","year":"2017","journal-title":"IEEE Commun. Mag."},{"issue":"5","key":"10.1016\/j.inffus.2026.104543_bib0025","doi-asserted-by":"crossref","first-page":"971","DOI":"10.1109\/JLT.1985.1074315","article-title":"Precision time domain reflectometry in optical fiber systems using a frequency modulated continuous wave ranging technique","volume":"3","author":"Uttam","year":"1985","journal-title":"J. Light. Technol."},{"key":"10.1016\/j.inffus.2026.104543_bib0026","unstructured":"C. Kong, The battle of LiDAR sensor technologies: FMCW vs. ToF, 2025. https:\/\/www.laserfocusworld.com\/test-measurement\/article\/55253453\/the-battle-of-lidar-sensor-technologies-fmcw-vs-tof."},{"key":"10.1016\/j.inffus.2026.104543_bib0027","series-title":"Drive-Thru Climate Tunnel: A Proposed Method to Study ADAS Performance in Adverse Weather","author":"Pao","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0028","doi-asserted-by":"crossref","first-page":"35174","DOI":"10.1109\/ACCESS.2025.3544350","article-title":"A study on the SLAM of automotive vehicles using bumper-mounted dual LiDAR","volume":"13","author":"Jang","year":"2025","journal-title":"IEEE Access"},{"key":"10.1016\/j.inffus.2026.104543_sbref0029","first-page":"194","article-title":"Assessing LiDAR sensor performance through automotive windshields with anti-reflective coatings","volume":"Spring 2025","author":"Song","year":"2025","journal-title":"Korean Soc. Mech. Eng."},{"issue":"2","key":"10.1016\/j.inffus.2026.104543_bib0030","doi-asserted-by":"crossref","first-page":"143","DOI":"10.1109\/JSEN.2001.936931","article-title":"An ultrasonic sensor for distance measurement in automotive applications","volume":"1","author":"Carullo","year":"2001","journal-title":"IEEE Sens. J."},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0031","doi-asserted-by":"crossref","first-page":"144","DOI":"10.54254\/2755-2721\/99\/20251773","article-title":"Applications of ultrasonic sensors: a review","volume":"99","author":"Wei","year":"2024","journal-title":"Appl. Comput. Eng."},{"key":"10.1016\/j.inffus.2026.104543_bib0032","series-title":"2020 IEEE Intelligent Vehicles Symposium (IV)","first-page":"2029","article-title":"Developments in modern GNSS and its impact on autonomous vehicle architectures","author":"Joubert","year":"2020"},{"issue":"3","key":"10.1016\/j.inffus.2026.104543_bib0033","first-page":"36","article-title":"Robust vehicular localization and map matching in urban environments through IMU, GNSS, and cellular signals","volume":"12","author":"Kassas","year":"2020","journal-title":"IEEE Intell. Transp. Syst. Mag."},{"issue":"8","key":"10.1016\/j.inffus.2026.104543_bib0034","doi-asserted-by":"crossref","first-page":"6469","DOI":"10.1109\/JIOT.2020.3043716","article-title":"Computing systems for autonomous driving: state of the art and challenges","volume":"8","author":"Liu","year":"2021","journal-title":"IEEE Internet Things J."},{"issue":"3","key":"10.1016\/j.inffus.2026.104543_bib0035","doi-asserted-by":"crossref","first-page":"355","DOI":"10.1177\/02783649241273554","article-title":"Dataset and benchmark: novel sensors for autonomous vehicle perception","volume":"44","author":"Carmichael","year":"2025","journal-title":"Int. J. Robot. Res."},{"key":"10.1016\/j.inffus.2026.104543_bib0036","doi-asserted-by":"crossref","first-page":"107710","DOI":"10.1109\/ACCESS.2021.3100472","article-title":"On 5G-V2X use cases and enabling technologies: a comprehensive survey","volume":"9","author":"Alalewi","year":"2021","journal-title":"IEEE Access"},{"key":"10.1016\/j.inffus.2026.104543_sbref0037","series-title":"DAIR-V2X: a large-scale dataset for vehicle-infrastructure cooperative 3D object detection","first-page":"21361","author":"Yu","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0038","series-title":"2022 OPJU International Technology Conference on Emerging Technologies for Sustainable Development (OTCON)","first-page":"1","article-title":"Application of microphone array in ADAS using model based development","author":"Oktey","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0039","doi-asserted-by":"crossref","DOI":"10.1016\/j.sna.2024.115586","article-title":"Pedestrian detection using a MEMS acoustic array mounted on a moving vehicle","volume":"376","author":"Izquierdo","year":"2024","journal-title":"Sens. Actuators A: Phys."},{"key":"10.1016\/j.inffus.2026.104543_bib0040","unstructured":"N.R. Shabtai, E. Tzirkel, Detecting the direction of emergency vehicle sirens with microphones, 2019, pp. 137. 10.25836\/sasp.2019.22."},{"key":"10.1016\/j.inffus.2026.104543_bib0041","doi-asserted-by":"crossref","first-page":"156465","DOI":"10.1109\/ACCESS.2021.3129150","article-title":"Object detection in thermal spectrum for advanced driver-assistance systems (ADAS)","volume":"9","author":"Farooq","year":"2021","journal-title":"IEEE Access"},{"key":"10.1016\/j.inffus.2026.104543_bib0042","series-title":"Digital Image Processing - Latest Advances and Applications","article-title":"Latest advancements in perception algorithms for ADAS and AV systems using infrared images and deep learning","author":"Srinivasan","year":"2023"},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0043","doi-asserted-by":"crossref","first-page":"154","DOI":"10.1109\/TPAMI.2020.3008413","article-title":"Event-based vision: a survey","volume":"44","author":"Gallego","year":"2022","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.inffus.2026.104543_bib0044","doi-asserted-by":"crossref","first-page":"51275","DOI":"10.1109\/ACCESS.2024.3386032","article-title":"Event cameras in automotive sensing: a review","volume":"12","author":"Shariff","year":"2024","journal-title":"IEEE Access"},{"key":"10.1016\/j.inffus.2026.104543_sbref0045","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"1506","article-title":"Gated2Depth: real-time dense lidar from gated images","author":"Gruber","year":"2019"},{"key":"10.1016\/j.inffus.2026.104543_sbref0046","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"2811","article-title":"Gated2Gated: self-supervised depth estimation from gated images","author":"Walia","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0047","doi-asserted-by":"crossref","first-page":"217","DOI":"10.1016\/j.robot.2018.11.023","article-title":"Extrinsic 6DoF calibration of a radar-LiDAR-camera system enhanced by radar cross section estimates evaluation","volume":"114","author":"Per\u0161i\u0107","year":"2019","journal-title":"Robot. Auton. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0048","doi-asserted-by":"crossref","first-page":"110417","DOI":"10.1109\/ACCESS.2023.3322229","article-title":"External extrinsic calibration of multi-modal imaging sensors: a review","volume":"11","author":"Liu","year":"2023","journal-title":"IEEE Access"},{"issue":"10","key":"10.1016\/j.inffus.2026.104543_bib0049","doi-asserted-by":"crossref","first-page":"17677","DOI":"10.1109\/TITS.2022.3155228","article-title":"Automatic extrinsic calibration method for LiDAR and camera sensor setups","volume":"23","author":"Beltr\u00e1n","year":"2022","journal-title":"IEEE Trans. Intell. Transp. Syst."},{"issue":"2","key":"10.1016\/j.inffus.2026.104543_bib0050","doi-asserted-by":"crossref","first-page":"4614","DOI":"10.1109\/LRA.2022.3151970","article-title":"Camera-IMU extrinsic calibration quality monitoring for autonomous ground vehicles","volume":"7","author":"Xiao","year":"2022","journal-title":"IEEE Robot. Autom. Lett."},{"key":"10.1016\/j.inffus.2026.104543_sbref0051","series-title":"The Future of Automotive Data Connectivity","author":"Arpe","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_sbref0052","series-title":"Automotive Ethernet: The in-Vehicle Networking of the Future","author":"Yeo","year":"2024"},{"key":"10.1016\/j.inffus.2026.104543_bib0053","series-title":"Getting Ready for Next-Generation E\/E Architecture with Zonal Compute","author":"Burkacky","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0054","series-title":"Software Architecture. ECSA 2022 Tracks and Workshops","first-page":"165","article-title":"Methodical approach for centralization evaluation of modern automotive e\/e architectures","author":"Mauser","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0055","series-title":"2013 IEEE Intelligent Vehicles Symposium (IV)","first-page":"763","article-title":"Towards a viable autonomous driving research platform","author":"Wei","year":"2013"},{"key":"10.1016\/j.inffus.2026.104543_bib0056","unstructured":"AUTOSAR adaptive platform, 2025. https:\/\/www.autosar.org\/standards\/adaptive-platform."},{"key":"10.1016\/j.inffus.2026.104543_bib0057","unstructured":"ROS: home, 2025. https:\/\/www.ros.org\/."},{"key":"10.1016\/j.inffus.2026.104543_bib0058","doi-asserted-by":"crossref","unstructured":"Z. Liu, H. Tang, A. Amini, X. Yang, H. Mao, D. Rus, S. Han, BEVFusion: multi-task multi-sensor fusion with unified bird\u2019s-eye view representation, (2024). 10.48550\/arXiv.2205.13542.","DOI":"10.1109\/ICRA48891.2023.10160968"},{"issue":"3","key":"10.1016\/j.inffus.2026.104543_bib0059","doi-asserted-by":"crossref","first-page":"2020","DOI":"10.1109\/TPAMI.2024.3515454","article-title":"BEVFormer: learning bird\u2019s-eye-view representation from LiDAR-camera via spatiotemporal transformers","volume":"47","author":"Li","year":"2025","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.inffus.2026.104543_bib0060","series-title":"Procedings of the British Machine Vision Conference 2013","first-page":"13.1","article-title":"Fast explicit diffusion for accelerated features in nonlinear scale spaces","author":"Alcantarilla","year":"2013"},{"key":"10.1016\/j.inffus.2026.104543_bib0061","series-title":"2011 International Conference on Computer Vision","first-page":"2564","article-title":"ORB: an efficient alternative to SIFT or SURF","author":"Rublee","year":"2011"},{"issue":"8","key":"10.1016\/j.inffus.2026.104543_bib0062","doi-asserted-by":"crossref","first-page":"3158","DOI":"10.1109\/TIP.2013.2259841","article-title":"Fast SIFT design for real-time visual feature extraction","volume":"22","author":"Chiu","year":"2013","journal-title":"IEEE Trans. Image Process."},{"issue":"3","key":"10.1016\/j.inffus.2026.104543_bib0063","doi-asserted-by":"crossref","first-page":"346","DOI":"10.1016\/j.cviu.2007.09.014","article-title":"Speeded-up robust features (SURF)","volume":"110","author":"Bay","year":"2008","journal-title":"Comput. Vis. Image Underst."},{"key":"10.1016\/j.inffus.2026.104543_bib0064","doi-asserted-by":"crossref","DOI":"10.7717\/peerj-cs.2415","article-title":"Comprehensive empirical evaluation of feature extractors in computer vision","volume":"10","author":"Isik","year":"2024","journal-title":"PeerJ Comput. Sci."},{"issue":"12","key":"10.1016\/j.inffus.2026.104543_bib0065","doi-asserted-by":"crossref","first-page":"10579","DOI":"10.1109\/TPAMI.2024.3444912","article-title":"Metric3D v2: a versatile monocular geometric foundation model for zero-shot metric depth and surface normal estimation","volume":"46","author":"Hu","year":"2024","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.inffus.2026.104543_sbref0066","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"9492","article-title":"Repurposing diffusion-based image generators for monocular depth estimation","author":"Ke","year":"2024"},{"key":"10.1016\/j.inffus.2026.104543_bib0067","series-title":"2012 IEEE\/RSJ International Conference on Intelligent Robots and Systems","first-page":"1644","article-title":"Evaluation of 3D feature descriptors for classification of surface geometries in point clouds","author":"Arbeiter","year":"2012"},{"issue":"16","key":"10.1016\/j.inffus.2026.104543_bib0068","doi-asserted-by":"crossref","first-page":"2127","DOI":"10.1016\/j.patrec.2012.07.006","article-title":"Hierarchical normal space sampling to speed up point cloud coarse matching","volume":"33","author":"Diez","year":"2012","journal-title":"Pattern Recognit. Lett."},{"issue":"11","key":"10.1016\/j.inffus.2026.104543_bib0069","doi-asserted-by":"crossref","first-page":"28099","DOI":"10.3390\/s151128099","article-title":"A review of LIDAR radiometric processing: from Ad Hoc intensity correction to rigorous radiometric calibration","volume":"15","author":"Kashani","year":"2015","journal-title":"Sensors"},{"key":"10.1016\/j.inffus.2026.104543_bib0070","series-title":"2016 IEEE 11th Conference on Industrial Electronics and Applications (ICIEA)","first-page":"2372","article-title":"Normalized radar cross section measurement for space objects","author":"Fu","year":"2016"},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0071","doi-asserted-by":"crossref","first-page":"130","DOI":"10.1109\/JOE.1987.1145241","article-title":"Two-dimensional normalization techniques","volume":"12","author":"Morgan","year":"1987","journal-title":"IEEE J. Ocean. Eng."},{"key":"10.1016\/j.inffus.2026.104543_bib0072","series-title":"2015 16th International Radar Symposium (IRS)","first-page":"174","article-title":"Clustering of high resolution automotive radar detections and subsequent feature extraction for classification of road users","author":"Schubert","year":"2015"},{"issue":"10","key":"10.1016\/j.inffus.2026.104543_bib0073","doi-asserted-by":"crossref","DOI":"10.3390\/s21103410","article-title":"Constraint-based hierarchical cluster selection in automotive radar data","volume":"21","author":"Malzer","year":"2021","journal-title":"Sensors"},{"key":"10.1016\/j.inffus.2026.104543_sbref0074","series-title":"Advances in Neural Information Processing Systems","article-title":"ImageNet classification with deep convolutional neural networks","volume":"Vol. 25","author":"Krizhevsky","year":"2012"},{"key":"10.1016\/j.inffus.2026.104543_bib0075","unstructured":"A. Dosovitskiy, L. Beyer, A. Kolesnikov, D. Weissenborn, X. Zhai, T. Unterthiner, M. Dehghani, M. Minderer, G. Heigold, S. Gelly, J. Uszkoreit, N. Houlsby, An image is worth 16x16 words: transformers for image recognition at scale, 2020. https:\/\/openreview.net\/forum?id=YicbFdNTTy."},{"key":"10.1016\/j.inffus.2026.104543_bib0076","doi-asserted-by":"crossref","unstructured":"K. He, X. Zhang, S. Ren, J. Sun, Deep residual learning for image recognition, 2016, pp. 770\u2013778. https:\/\/openaccess.thecvf.com\/content_cvpr_2016\/html\/He_Deep_Residual_Learning_CVPR_2016_paper.html.","DOI":"10.1109\/CVPR.2016.90"},{"key":"10.1016\/j.inffus.2026.104543_sbref0077","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"2117","article-title":"Feature pyramid networks for object detection","author":"Lin","year":"2017"},{"key":"10.1016\/j.inffus.2026.104543_sbref0078","series-title":"Proceedings of the European Conference on Computer Vision (ECCV)","first-page":"3","article-title":"Group normalization","author":"Wu","year":"2018"},{"key":"10.1016\/j.inffus.2026.104543_sbref0079","series-title":"Deformable convolutional networks","first-page":"764","author":"Dai","year":"2017"},{"key":"10.1016\/j.inffus.2026.104543_bib0080","unstructured":"X. Zhu, W. Su, L. Lu, B. Li, X. Wang, J. Dai, Deformable DETR: deformable transformers for end-to-end object detection, (2021). 10.48550\/arXiv.2010.04159."},{"key":"10.1016\/j.inffus.2026.104543_sbref0081","series-title":"Proceedings of the 33rd International Conference on Machine Learning","first-page":"2990","article-title":"Group equivariant convolutional networks","author":"Cohen","year":"2016"},{"issue":"20","key":"10.1016\/j.inffus.2026.104543_bib0082","doi-asserted-by":"crossref","first-page":"24551","DOI":"10.1007\/s10489-023-04747-6","article-title":"GET: group equivariant transformer for person detection of overhead fisheye images","volume":"53","author":"Chen","year":"2023","journal-title":"Appl. Intell."},{"key":"10.1016\/j.inffus.2026.104543_bib0083","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"PointNet: deep learning on point sets for 3D classification and segmentation","author":"Qi","year":"2017"},{"key":"10.1016\/j.inffus.2026.104543_sbref0084","series-title":"Advances in Neural Information Processing Systems","article-title":"PointNet++: deep hierarchical feature learning on point sets in a metric space","volume":"30","author":"Qi","year":"2017"},{"key":"10.1016\/j.inffus.2026.104543_sbref0085","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"6411","article-title":"KPConv: flexible and deformable convolution for point clouds","author":"Thomas","year":"2019"},{"key":"10.1016\/j.inffus.2026.104543_bib0086","doi-asserted-by":"crossref","unstructured":"W. Shi, R. Rajkumar, Point-GNN: graph neural network for 3D object detection in a point cloud, 2020, pp. 1711\u20131719. https:\/\/openaccess.thecvf.com\/content_CVPR_2020\/html\/Shi_Point-GNN_Graph_Neural_Network_for_3D_Object_Detection_in_a_CVPR_2020_paper.html.","DOI":"10.1109\/CVPR42600.2020.00178"},{"key":"10.1016\/j.inffus.2026.104543_bib0087","doi-asserted-by":"crossref","unstructured":"F. Fent, P. Bauerschmidt, M. Lienkamp, RadarGNN: transformation invariant graph neural network for radar-based perception, 2023, pp. 182\u2013191. https:\/\/openaccess.thecvf.com\/content\/CVPR2023W\/WAD\/html\/Fent_RadarGNN_Transformation_Invariant_Graph_Neural_Network_for_Radar-Based_Perception_CVPRW_2023_paper.html.","DOI":"10.1109\/CVPRW59228.2023.00023"},{"key":"10.1016\/j.inffus.2026.104543_bib0088","doi-asserted-by":"crossref","unstructured":"Y. Zhou, O. Tuzel, VoxelNet: end-to-end learning for point cloud based 3d object detection, 2018, pp. 4490\u20134499. https:\/\/openaccess.thecvf.com\/content_cvpr_2018\/html\/Zhou_VoxelNet_End-to-End_Learning_CVPR_2018_paper.html.","DOI":"10.1109\/CVPR.2018.00472"},{"key":"10.1016\/j.inffus.2026.104543_bib0089","doi-asserted-by":"crossref","unstructured":"A.H. Lang, S. Vora, H. Caesar, L. Zhou, J. Yang, O. Beijbom, PointPillars: fast encoders for object detection from point clouds, 2019, pp. 12697\u201312705. https:\/\/openaccess.thecvf.com\/content_CVPR_2019\/html\/Lang_PointPillars_Fast_Encoders_for_Object_Detection_From_Point_Clouds_CVPR_2019_paper.html.","DOI":"10.1109\/CVPR.2019.01298"},{"issue":"7","key":"10.1016\/j.inffus.2026.104543_bib0090","doi-asserted-by":"crossref","DOI":"10.3390\/app14072781","article-title":"Sparsity-robust feature fusion for vulnerable road-user detection with 4D radar","volume":"14","author":"Ruddat","year":"2024","journal-title":"Appl. Sci."},{"issue":"17","key":"10.1016\/j.inffus.2026.104543_bib0091","doi-asserted-by":"crossref","first-page":"4393","DOI":"10.3390\/rs14174393","article-title":"Generalized LiDAR intensity normalization and its positive impact on geometric and learning-based lane marking detection","volume":"14","author":"Cheng","year":"2022","journal-title":"Remote Sens."},{"key":"10.1016\/j.inffus.2026.104543_bib0092","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"16259","article-title":"Point transformer","author":"Zhao","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0093","doi-asserted-by":"crossref","first-page":"33330","DOI":"10.52202\/068431-2415","article-title":"Point transformer v2: grouped vector attention and partition-based pooling","volume":"35","author":"Wu","year":"2022","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0094","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition","first-page":"4840","article-title":"Point transformer v3: simpler faster stronger","author":"Wu","year":"2024"},{"key":"10.1016\/j.inffus.2026.104543_sbref0095","series-title":"International Conference on Learning Representations","article-title":"Tent: fully test-time adaptation by entropy minimization","author":"Wang","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0096","doi-asserted-by":"crossref","first-page":"1106","DOI":"10.1007\/s11263-024-02213-5","article-title":"In search of lost online test-time adaptation: a survey","volume":"133","author":"Wang","year":"2025","journal-title":"Int. J. Comput. Vis."},{"key":"10.1016\/j.inffus.2026.104543_bib0097","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"7201","article-title":"Continual test-time domain adaptation","author":"Wang","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0098","series-title":"Proceedings of the 39th International Conference on Machine Learning","first-page":"16888","article-title":"Efficient test-time model adaptation without forgetting","volume":"162","author":"Niu","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_sbref0099","series-title":"Proceedings of the 37th International Conference on Machine Learning","first-page":"1597","article-title":"A simple framework for contrastive learning of visual representations","author":"Chen","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0100","series-title":"2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"9726","article-title":"Momentum contrast for unsupervised visual representation learning","author":"He","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_sbref0101","series-title":"Advances in Neural Information Processing Systems","first-page":"21271","article-title":"Bootstrap your own latent - a new approach to self-supervised learning","volume":"Vol. 33","author":"Grill","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0102","doi-asserted-by":"crossref","unstructured":"M. Caron, H. Touvron, I. Misra, H. J\u00e9gou, J. Mairal, P. Bojanowski, A. Joulin, Emerging properties in self-supervised vision transformers, 2021, pp. 9650\u20139660. https:\/\/openaccess.thecvf.com\/content\/ICCV2021\/html\/Caron_Emerging_Properties_in_Self-Supervised_Vision_Transformers_ICCV_2021_paper.","DOI":"10.1109\/ICCV48922.2021.00951"},{"key":"10.1016\/j.inffus.2026.104543_bib0103","series-title":"2024 International Conference on 3D Vision (3DV)","first-page":"559","article-title":"BEVContrast: self-supervision in BEV space for automotive lidar point clouds","author":"Sautier","year":"2024"},{"issue":"12","key":"10.1016\/j.inffus.2026.104543_bib0104","doi-asserted-by":"crossref","first-page":"22167","DOI":"10.1109\/JIOT.2024.3379471","article-title":"BEVSOC: self-supervised contrastive learning for calibration-free BEV 3-D object detection","volume":"11","author":"Chen","year":"2024","journal-title":"IEEE Internet Things J."},{"key":"10.1016\/j.inffus.2026.104543_bib0105","series-title":"Computer Vision - ECCV 2020","first-page":"194","article-title":"Lift, splat, shoot: encoding images from arbitrary camera rigs by implicitly unprojecting to 3D","author":"Philion","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0106","doi-asserted-by":"crossref","first-page":"11713","DOI":"10.1109\/LRA.2025.3611145","article-title":"Efficient multi-camera tokenization with triplanes for end-to-end driving","volume":"10","author":"Ivanovic","year":"2025","journal-title":"IEEE Robot. Autom. Lett."},{"key":"10.1016\/j.inffus.2026.104543_bib0107","unstructured":"J. Huang, G. Huang, BEVDet4D: exploit temporal cues in multi-camera 3D object detection, (2022). 10.48550\/arXiv.2203.17054."},{"issue":"12","key":"10.1016\/j.inffus.2026.104543_bib0108","doi-asserted-by":"crossref","first-page":"8665","DOI":"10.1109\/TPAMI.2024.3414835","article-title":"Fast-BEV: a fast and strong bird\u2019s-eye view perception baseline","volume":"46","author":"Li","year":"2024","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.inffus.2026.104543_bib0109","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"8515","article-title":"Towards viewpoint robustness in bird\u2019s eye view segmentation","author":"Klinghoffer","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0110","unstructured":"X. Liu, H. Shen, Benchmarking multi-view BEV object detection with mixed pinhole and fisheye cameras, arXiv: 2603.27818(2026)."},{"key":"10.1016\/j.inffus.2026.104543_sbref0111","series-title":"Proceedings of the 5th Conference on Robot Learning","first-page":"180","article-title":"DETR3D: 3D object detection from multi-view images via 3D-to-2D queries","author":"Wang","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0112","series-title":"Computer Vision - ECCV 2022","first-page":"531","article-title":"PETR: position embedding transformation for multi-view 3D object detection","author":"Liu","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0113","doi-asserted-by":"crossref","unstructured":"L. Peng, Z. Chen, Z. Fu, P. Liang, E. Cheng, BEVSegFormer: bird\u2019s eye view semantic segmentation from arbitrary camera rigs, 2023, pp. 5935\u20135943. https:\/\/openaccess.thecvf.com\/content\/WACV2023\/html\/Peng_BEVSegFormer_Birds_Eye_View_Semantic_Segmentation_From_Arbitrary_Camera_Rigs_WACV_2023_paper.html.","DOI":"10.1109\/WACV56688.2023.00588"},{"key":"10.1016\/j.inffus.2026.104543_bib0114","series-title":"International Conference on Machine Learning","first-page":"4651","article-title":"Perceiver: general perception with iterative attention","author":"Jaegle","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0115","series-title":"The Thirty-ninth Annual Conference on Neural Information Processing Systems","article-title":"Rig3R: rig-aware conditioning and discovery for 3D reconstruction","author":"Li","year":"2025"},{"key":"10.1016\/j.inffus.2026.104543_bib0116","unstructured":"J. Yang, Z. Chen, Y. You, Y. Wang, Y. Li, Y. Chen, B. Li, B. Ivanovic, M. Pavone, Y. Wang, Towards efficient and effective multi-camera encoding for end-to-end driving, arXiv: 2512.10947(2025). Flex scene encoder with compact learned tokens for LLM-based driving policies."},{"key":"10.1016\/j.inffus.2026.104543_bib0117","doi-asserted-by":"crossref","first-page":"16344","DOI":"10.52202\/068431-1189","article-title":"Flashattention: fast and memory-efficient exact attention with io-awareness","volume":"35","author":"Dao","year":"2022","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0118","series-title":"European Conference on Computer Vision","first-page":"348","article-title":"Multimae: multi-modal multi-task masked autoencoders","author":"Bachmann","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_bib0119","doi-asserted-by":"crossref","unstructured":"R. Girdhar, A. El-Nouby, Z. Liu, M. Singh, K.V. Alwala, A. Joulin, I. Misra, ImageBind: one embedding space to bind them all, 2023, pp. 15180\u201315190. https:\/\/openaccess.thecvf.com\/content\/CVPR2023\/html\/Girdhar_ImageBind_One_Embedding_Space_To_Bind_Them_All_CVPR_2023_paper.html.","DOI":"10.1109\/CVPR52729.2023.01457"},{"key":"10.1016\/j.inffus.2026.104543_bib0120","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"18268","article-title":"Cross modal transformer: towards fast and robust 3d object detection","author":"Yan","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0121","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"8721","article-title":"Metabev: solving sensor failures for 3d detection and map segmentation","author":"Ge","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0122","series-title":"2024 IEEE Intelligent Vehicles Symposium (IV)","first-page":"2776","article-title":"Unibev: multi-modal 3d object detection with uniform bev encoders for robustness against missing sensor modalities","author":"Wang","year":"2024"},{"key":"10.1016\/j.inffus.2026.104543_bib0123","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition","first-page":"6720","article-title":"Resilient sensor fusion under adverse sensor failures via multi-modal expert fusion","author":"Park","year":"2025"},{"key":"10.1016\/j.inffus.2026.104543_bib0124","unstructured":"H. Jiang, W. Meng, H. Zhu, Q. Zhang, J. Yin, Multi-camera calibration free BEV representation for 3D object detection, (2022). 10.48550\/arXiv.2210.17252."},{"key":"10.1016\/j.inffus.2026.104543_bib0125","doi-asserted-by":"crossref","unstructured":"S. Wang, V. Leroy, Y. Cabon, B. Chidlovskii, J. Revaud, DUSt3R: geometric 3D vision made easy, 2024, pp. 20697\u201320709. https:\/\/openaccess.thecvf.com\/content\/CVPR2024\/html\/Wang_DUSt3R_Geometric_3D_Vision_Made_Easy_CVPR_2024_paper.html.","DOI":"10.1109\/CVPR52733.2024.01956"},{"key":"10.1016\/j.inffus.2026.104543_bib0126","doi-asserted-by":"crossref","unstructured":"J. Yang, A. Sax, K.J. Liang, M. Henaff, H. Tang, A. Cao, J. Chai, F. Meier, M. Feiszli, Fast3R: towards 3D reconstruction of 1000+ images in one forward pass, 2025, pp. 21924\u201321935. https:\/\/openaccess.thecvf.com\/content\/CVPR2025\/html\/Yang_Fast3R_Towards_3D_Reconstruction_of_1000_Images_in_One_Forward_CVPR_2025_paper.html.","DOI":"10.1109\/CVPR52734.2025.02042"},{"key":"10.1016\/j.inffus.2026.104543_bib0127","doi-asserted-by":"crossref","unstructured":"J. Wang, M. Chen, N. Karaev, A. Vedaldi, C. Rupprecht, D. Novotny, VGGT: visual geometry grounded transformer, 2025, pp. 5294\u20135306. https:\/\/openaccess.thecvf.com\/content\/CVPR2025\/html\/Wang_VGGT_Visual_Geometry_Grounded_Transformer_CVPR_2025_paper.html.","DOI":"10.1109\/CVPR52734.2025.00499"},{"key":"10.1016\/j.inffus.2026.104543_bib0128","series-title":"2020 IEEE\/RSJ International Conference on Intelligent Robots and Systems (IROS)","first-page":"10386","article-title":"CLOCs: camera-LiDAR object candidates fusion for 3D object detection","author":"Pang","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_sbref0129","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"17545","article-title":"SparseFusion: fusing multi-modal sparse representations for multi-sensor 3D object detection","author":"Xie","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0130","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"18067","article-title":"ObjectFusion: multi-modal 3D object detection with object-centric fusion","author":"Cai","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0131","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"1090","article-title":"TransFusion: robust LiDAR-camera fusion for 3D object detection with transformers","author":"Bai","year":"2022"},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0132","doi-asserted-by":"crossref","first-page":"28","DOI":"10.1016\/j.inffus.2011.08.001","article-title":"Multisensor data fusion: a review of the state-of-the-art","volume":"14","author":"Khaleghi","year":"2013","journal-title":"Inf. Fusion"},{"key":"10.1016\/j.inffus.2026.104543_bib0133","doi-asserted-by":"crossref","unstructured":"Z. Qin, J. Chen, C. Chen, X. Chen, X. Li, UniFusion: unified multi-view fusion transformer for spatial-temporal representation in bird\u2019s-eye-view, 2023, pp. 8690\u20138699. https:\/\/openaccess.thecvf.com\/content\/ICCV2023\/html\/Qin_UniFusion_Unified_Multi-View_Fusion_Transformer_for_Spatial-Temporal_Representation_in_Birds-Eye-View_ICCV_2023_paper.html.","DOI":"10.1109\/ICCV51070.2023.00798"},{"key":"10.1016\/j.inffus.2026.104543_bib0134","doi-asserted-by":"crossref","unstructured":"L. Zheng, J. Liu, R. Guan, L. Yang, S. Lu, Y. Li, X. Bai, J. Bai, Z. Ma, H.-L. Shen, X. Zhu, Doracamom: joint 3D detection and occupancy prediction with multi-view 4D radars and cameras for omnidirectional perception, (2025). 10.48550\/arXiv.2501.15394.","DOI":"10.1109\/TCSVT.2026.3657111"},{"key":"10.1016\/j.inffus.2026.104543_sbref0135","series-title":"International Conference on Learning Representations","article-title":"Time will tell: new outlooks and a baseline for temporal multi-view 3D object detection","author":"Park","year":"2023"},{"issue":"7","key":"10.1016\/j.inffus.2026.104543_bib0136","doi-asserted-by":"crossref","first-page":"6544","DOI":"10.1109\/LRA.2024.3401172","article-title":"Exploring recurrent long-term temporal fusion for multi-view 3d perception","volume":"9","author":"Han","year":"2024","journal-title":"IEEE Robot. Autom. Lett."},{"key":"10.1016\/j.inffus.2026.104543_bib0137","series-title":"Computer Vision \u2013 ECCV 2024","first-page":"131","article-title":"RecurrentBEV: a long-term temporal fusion framework for multi-view 3D detection","author":"Chang","year":"2024"},{"key":"10.1016\/j.inffus.2026.104543_bib0138","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"3239","article-title":"PETRv2: a unified framework for 3D perception from multi-camera images","author":"Liu","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0139","unstructured":"X. Lin, T. Lin, Z. Pei, L. Huang, Z. Su, Sparse4D: multi-view 3D object detection with sparse spatial-temporal fusion, (2023). 10.48550\/arXiv.2211.10581."},{"key":"10.1016\/j.inffus.2026.104543_bib0140","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"3621","article-title":"Exploring object-centric temporal modeling for efficient multi-view 3D object detection","author":"Wang","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0141","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"18580","article-title":"SparseBEV: high-performance sparse 3D object detection from multi-camera videos","author":"Liu","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0142","unstructured":"M. Fan, Y. Zuo, P. Blaes, H. Montgomery, S. Das, Robust sensor fusion against on-vehicle sensor staleness, (2025). 10.48550\/arXiv.2506.05780."},{"key":"10.1016\/j.inffus.2026.104543_bib0143","unstructured":"D. Ha, J. Schmidhuber, World models, (2018). 10.48550\/arXiv.1803.10122."},{"key":"10.1016\/j.inffus.2026.104543_bib0144","first-page":"1","article-title":"World models for autonomous driving: an initial survey","author":"Guan","year":"2024","journal-title":"IEEE Trans. Intell. Veh."},{"key":"10.1016\/j.inffus.2026.104543_bib0145","doi-asserted-by":"crossref","unstructured":"D. Hafner, J. Pasukonis, J. Ba, T. Lillicrap, Mastering diverse domains through world models, (2024). 10.48550\/arXiv.2301.04104.","DOI":"10.1038\/s41586-025-08744-2"},{"key":"10.1016\/j.inffus.2026.104543_bib0146","unstructured":"F. Jia, W. Mao, Y. Liu, Y. Zhao, Y. Wen, C. Zhang, X. Zhang, T. Wang, ADriver-I: a general world model for autonomous driving, (2023). 10.48550\/arXiv.2311.13549."},{"key":"10.1016\/j.inffus.2026.104543_bib0147","unstructured":"A. Hu, L. Russell, H. Yeo, Z. Murez, G. Fedoseev, A. Kendall, J. Shotton, G. Corrado, GAIA-1: a generative world model for autonomous driving, (2023). 10.48550\/arXiv.2309.17080."},{"key":"10.1016\/j.inffus.2026.104543_bib0148","doi-asserted-by":"crossref","unstructured":"D. Bogdoll, Y. Yang, T. Joseph, M. Yazgan, J.M. Z\u00f6llner, MUVO: a multimodal generative world model for autonomous driving with geometric representations, (2025). 10.48550\/arXiv.2311.11762.","DOI":"10.1109\/IV64158.2025.11097718"},{"key":"10.1016\/j.inffus.2026.104543_bib0149","series-title":"Computer Vision - ECCV 2024","first-page":"55","article-title":"Drivedreamer: towards real-world-drive world models for autonomous driving","author":"Wang","year":"2025"},{"key":"10.1016\/j.inffus.2026.104543_bib0150","series-title":"Computer Vision - ECCV 2024","first-page":"142","article-title":"Think2Drive: efficient reinforcement learning by thinking with latent world model for autonomous driving (in CARLA-V2)","author":"Li","year":"2025"},{"key":"10.1016\/j.inffus.2026.104543_bib0151","series-title":"Computer Vision - ECCV 2014","first-page":"740","article-title":"Microsoft COCO: common objects in context","author":"Lin","year":"2014"},{"key":"10.1016\/j.inffus.2026.104543_bib0152","doi-asserted-by":"crossref","unstructured":"M. Cordts, M. Omran, S. Ramos, T. Rehfeld, M. Enzweiler, R. Benenson, U. Franke, S. Roth, B. Schiele, The cityscapes dataset for semantic urban scene understanding, 2016, pp. 3213\u20133223. https:\/\/openaccess.thecvf.com\/content_cvpr_2016\/html\/Cordts_The_Cityscapes_Dataset_CVPR_2016_paper.html.","DOI":"10.1109\/CVPR.2016.350"},{"key":"10.1016\/j.inffus.2026.104543_sbref0153","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"11621","article-title":"nuScenes: a multimodal dataset for autonomous driving","author":"Caesar","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0154","series-title":"2021IEEE International Conference on Robotics and Automation (ICRA)","first-page":"6732","article-title":"Lightweight semantic mesh mapping for autonomous vehicles","author":"Herb","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0155","doi-asserted-by":"crossref","unstructured":"W. Tong, C. Sima, T. Wang, L. Chen, S. Wu, H. Deng, Y. Gu, L. Lu, P. Luo, D. Lin, H. Li, Scene as Occupancy, 2023, pp. 8406\u20138415. https:\/\/openaccess.thecvf.com\/content\/ICCV2023\/html\/Tong_Scene_as_Occupancy_ICCV_2023_paper.html.","DOI":"10.1109\/ICCV51070.2023.00772"},{"key":"10.1016\/j.inffus.2026.104543_sbref0156","series-title":"SurroundOcc: multi-camera 3D occupancy prediction for autonomous driving","first-page":"21729","author":"Wei","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_sbref0157","series-title":"OpenOccupancy: a large scale benchmark for surrounding semantic occupancy perception","first-page":"17850","author":"Wang","year":"2023"},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0158","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1109\/TPAMI.2021.3137605","article-title":"A comprehensive survey of scene graphs: generation and application","volume":"45","author":"Chang","year":"2023","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0159","first-page":"91","article-title":"Mapping for autonomous driving: opportunities and challenges","volume":"13","author":"Wong","year":"2021","journal-title":"IEEE Intell. Transp. Syst. Mag."},{"issue":"7","key":"10.1016\/j.inffus.2026.104543_bib0160","doi-asserted-by":"crossref","first-page":"11814","DOI":"10.1109\/TNNLS.2024.3495045","article-title":"Grid-centric traffic scenario perception for autonomous driving: a comprehensive review","volume":"36","author":"Shi","year":"2025","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0161","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition","first-page":"11784","article-title":"Center-based 3d object detection and tracking","author":"Yin","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0162","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision","first-page":"913","article-title":"Fcos3d: fully convolutional one-stage monocular 3d object detection","author":"Wang","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0163","series-title":"Computer Vision - ECCV 2020","first-page":"213","article-title":"End-to-end object detection with transformers","author":"Carion","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0164","unstructured":"A. Jaegle, S. Borgeaud, J.-B. Alayrac, C. Doersch, C. Ionescu, D. Ding, S. Koppula, D. Zoran, A. Brock, E. Shelhamer, et al., Perceiver IO: a general architecture for structured inputs & outputs, arXiv: 2107.14795(2021)."},{"key":"10.1016\/j.inffus.2026.104543_bib0165","unstructured":"O.A. van den, Y. Li, O. Vinyals, Representation learning with contrastive predictive coding, (2019). 10.48550\/arXiv.1807.03748."},{"key":"10.1016\/j.inffus.2026.104543_bib0166","series-title":"Computer Vision - ECCV 2020","first-page":"776","article-title":"Contrastive multiview coding","author":"Tian","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_sbref0167","series-title":"Proceedings of the 38th International Conference on Machine Learning","first-page":"8748","article-title":"Learning transferable visual models from natural language supervision","author":"Radford","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0168","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"11975","article-title":"Sigmoid loss for language image pre-training","author":"Zhai","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0169","unstructured":"M. Tschannen, A. Gritsenko, X. Wang, M.F. Naeem, I. Alabdulmohsin, N. Parthasarathy, T. Evans, L. Beyer, Y. Xia, B. Mustafa, O. Henaff, J. Harmsen, A. Steiner, X. Zhai, SigLIP 2: multilingual vision-language encoders with improved semantic understanding, localization, and dense features, (2025). 10.48550\/arXiv.2502.14786."},{"key":"10.1016\/j.inffus.2026.104543_bib0170","unstructured":"M. Oquab, T. Darcet, T. Moutakanni, H. Vo, M. Szafraniec, V. Khalidov, P. Fernandez, D. Haziza, F. Massa, A. El-Nouby, M. Assran, N. Ballas, W. Galuba, R. Howes, P.-Y. Huang, S.-W. Li, I. Misra, M. Rabbat, V. Sharma, G. Synnaeve, H. Xu, H. Jegou, J. Mairal, P. Labatut, A. Joulin, P. Bojanowski, DINOv2: learning robust visual features without supervision, (2024). 10.48550\/arXiv.2304.07193."},{"key":"10.1016\/j.inffus.2026.104543_bib0171","unstructured":"O. Sim\u00e9oni, H.V. Vo, M. Seitzer, F. Baldassarre, M. Oquab, C. Jose, V. Khalidov, M. Szafraniec, S. Yi, M. Ramamonjisoa, F. Massa, D. Haziza, L. Wehrstedt, J. Wang, T. Darcet, T. Moutakanni, L. Sentana, C. Roberts, A. Vedaldi, J. Tolan, J. Brandt, C. Couprie, J. Mairal, H. J\u00e9gou, P. Labatut, P. Bojanowski, DINOv3, (2025). 10.48550\/arXiv.2508.10104."},{"key":"10.1016\/j.inffus.2026.104543_sbref0172","series-title":"Proceedings of the 38th International Conference on Machine Learning","first-page":"12310","article-title":"Barlow twins: self-supervised learning via redundancy reduction","author":"Zbontar","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0173","unstructured":"A. Bardes, J. Ponce, Y. LeCun, VICReg: variance-invariance-covariance regularization for self-supervised learning, (2022). 10.48550\/arXiv.2105.04906."},{"key":"10.1016\/j.inffus.2026.104543_bib0174","unstructured":"D.P. Kingma, M. Welling, Auto-encoding variational bayes(2013). https:\/\/openreview.net\/forum?id=33X9fd2-9FyZd."},{"key":"10.1016\/j.inffus.2026.104543_bib0175","unstructured":"I. Higgins, L. Matthey, A. Pal, C. Burgess, X. Glorot, M. Botvinick, S. Mohamed, A. Lerchner, beta-VAE: learning basic visual concepts with a constrained variational framework, 2017. https:\/\/openreview.net\/forum?id=Sy2fzU9gl."},{"key":"10.1016\/j.inffus.2026.104543_sbref0176","series-title":"Proceedings of the 35th International Conference on Machine Learning","first-page":"2649","article-title":"Disentangling by factorising","author":"Kim","year":"2018"},{"key":"10.1016\/j.inffus.2026.104543_sbref0177","series-title":"Proceedings of the 36th International Conference on Machine Learning","first-page":"4114","article-title":"Challenging common assumptions in the unsupervised learning of disentangled representations","author":"Locatello","year":"2019"},{"issue":"12","key":"10.1016\/j.inffus.2026.104543_bib0178","doi-asserted-by":"crossref","first-page":"9677","DOI":"10.1109\/TPAMI.2024.3420937","article-title":"Disentangled representation learning","volume":"46","author":"Wang","year":"2024","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.inffus.2026.104543_sbref0179","series-title":"Advances in Neural Information Processing Systems","article-title":"Generative adversarial nets","volume":"Vol. 27","author":"Goodfellow","year":"2014"},{"key":"10.1016\/j.inffus.2026.104543_sbref0180","series-title":"Advances in Neural Information Processing Systems","article-title":"InfoGAN: interpretable representation learning by information maximizing generative adversarial nets","volume":"vol. 29","author":"Chen","year":"2016"},{"key":"10.1016\/j.inffus.2026.104543_bib0181","doi-asserted-by":"crossref","unstructured":"R. Rombach, A. Blattmann, D. Lorenz, P. Esser, B. Ommer, High-resolution image synthesis with latent diffusion models, 2022, pp. 10684\u201310695. https:\/\/openaccess.thecvf.com\/content\/CVPR2022\/html\/Rombach_High-Resolution_Image_Synthesis_With_Latent_Diffusion_Models_CVPR_2022_paper.","DOI":"10.1109\/CVPR52688.2022.01042"},{"key":"10.1016\/j.inffus.2026.104543_bib0182","series-title":"Computer Vision \u2013 ECCV 2020","first-page":"447","article-title":"Model-based occlusion disentanglement for image-to-image translation","author":"Pizzati","year":"2020"},{"key":"10.1016\/j.inffus.2026.104543_bib0183","unstructured":"M. Arjovsky, L. Bottou, I. Gulrajani, D. Lopez-Paz, Invariant risk minimization, (2020). 10.48550\/arXiv.1907.02893."},{"key":"10.1016\/j.inffus.2026.104543_sbref0184","series-title":"Proceedings of the 38th International Conference on Machine Learning","first-page":"2189","article-title":"Environment inference for invariant learning","author":"Creager","year":"2021"},{"key":"10.1016\/j.inffus.2026.104543_bib0185","unstructured":"I. Gulrajani, D. Lopez-Paz, In search of lost domain generalization, 2020. https:\/\/openreview.net\/forum?id=lQdXeXDoWtI."},{"issue":"6088","key":"10.1016\/j.inffus.2026.104543_bib0186","doi-asserted-by":"crossref","first-page":"533","DOI":"10.1038\/323533a0","article-title":"Learning representations by back-propagating errors","volume":"323","author":"Rumelhart","year":"1986","journal-title":"Nature"},{"issue":"7","key":"10.1016\/j.inffus.2026.104543_bib0187","doi-asserted-by":"crossref","first-page":"5753","DOI":"10.1109\/TCSVT.2024.3366664","article-title":"Toward robust LiDAR-camera fusion in BEV space via mutual deformable attention and temporal aggregation","volume":"34","author":"Wang","year":"2024","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.inffus.2026.104543_bib0188","article-title":"What uncertainties do we need in Bayesian deep learning for computer vision?","volume":"30","author":"Kendall","year":"2017","journal-title":"Adv. Neural Inf. Process. Syst."},{"issue":"2","key":"10.1016\/j.inffus.2026.104543_bib0189","doi-asserted-by":"crossref","first-page":"3153","DOI":"10.1109\/LRA.2020.2974682","article-title":"A general framework for uncertainty estimation in deep learning","volume":"5","author":"Loquercio","year":"2020","journal-title":"IEEE Robot. Autom. Lett."},{"key":"10.1016\/j.inffus.2026.104543_bib0190","series-title":"Proceedings of the 27th International Conference on Intelligent User Interfaces","first-page":"173","article-title":"Deep learning uncertainty in machine teaching","author":"Sanchez","year":"2022"},{"key":"10.1016\/j.inffus.2026.104543_sbref0191","series-title":"Proceedings of the 33rd International Conference on Machine Learning","first-page":"1050","article-title":"Dropout as a Bayesian approximation: representing model uncertainty in deep learning","author":"Gal","year":"2016"},{"key":"10.1016\/j.inffus.2026.104543_bib0192","article-title":"Can you trust your model\u2019s uncertainty? Evaluating predictive uncertainty under dataset shift","volume":"32","author":"Ovadia","year":"2019","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0193","first-page":"21464","article-title":"Energy-based out-of-distribution detection","volume":"33","author":"Liu","year":"2020","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"10.1016\/j.inffus.2026.104543_bib0194","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition","first-page":"24384","article-title":"Deep deterministic uncertainty: a new simple baseline","author":"Mukhoti","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0195","doi-asserted-by":"crossref","first-page":"180-1","DOI":"10.2352\/ISSN.2470-1173.2021.17.AVM-180","article-title":"Data driven degradation of automotive sensors and effect analysis","volume":"33","author":"Fleck","year":"2021","journal-title":"Electron. Imaging"},{"key":"10.1016\/j.inffus.2026.104543_bib0196","series-title":"2023 IEEE Intelligent Vehicles Symposium (IV)","first-page":"1","article-title":"Survey on liDAR perception in adverse weather conditions","author":"Dreissig","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0197","doi-asserted-by":"crossref","first-page":"2857","DOI":"10.1109\/OJVT.2025.3621862","article-title":"Exploring sensor impact and architectural robustness in adverse weather on BEV perception","volume":"6","author":"Kumar","year":"2025","journal-title":"IEEE Open J. Veh. Technol."},{"key":"10.1016\/j.inffus.2026.104543_bib0198","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"3354","article-title":"Are we ready for autonomous driving? the KITTI vision benchmark suite","author":"Geiger","year":"2012"},{"key":"10.1016\/j.inffus.2026.104543_bib0199","doi-asserted-by":"crossref","unstructured":"P. Sun, H. Kretzschmar, X. Dotiwalla, A. Chouard, V. Patnaik, P. Tsui, J. Guo, Y. Zhou, Y. Chai, B. Caine, V. Vasudevan, W. Han, J. Ngiam, H. Zhao, A. Timofeev, S. Ettinger, M. Krivokon, A. Gao, A. Joshi, Y. Zhang, J. Shlens, Z. Chen, D. Anguelov, Scalability in perception for autonomous driving: waymo open dataset, 2020, pp. 2446\u20132454. https:\/\/openaccess.thecvf.com\/content_CVPR_2020\/html\/Sun_Scalability_in_Perception_for_Autonomous_Driving_Waymo_Open_Dataset_CVPR_2020_paper.html.","DOI":"10.1109\/CVPR42600.2020.00252"},{"issue":"11","key":"10.1016\/j.inffus.2026.104543_bib0200","doi-asserted-by":"crossref","first-page":"7138","DOI":"10.1109\/TIV.2024.3394735","article-title":"A survey on autonomous driving datasets: statistics, annotation quality, and a future outlook","volume":"9","author":"Liu","year":"2024","journal-title":"IEEE Trans. Intell. Veh."},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0201","doi-asserted-by":"crossref","first-page":"1847","DOI":"10.1109\/TIV.2023.3331024","article-title":"Synthetic datasets for autonomous driving: a survey","volume":"9","author":"Song","year":"2024","journal-title":"IEEE Trans. Intell. Veh."},{"key":"10.1016\/j.inffus.2026.104543_bib0202","unstructured":"J. Geyer, Y. Kassahun, M. Mahmudi, X. Ricou, R. Durgesh, A.S. Chung, L. Hauswald, V.H. Pham, M. M\u00fchlegg, S. Dorn, T. Fernandez, M. J\u00e4nicke, S. Mirashi, C. Savani, M. Sturm, O. Vorobiov, M. Oelker, S. Garreis, P. Schuberth, A2D2: audi autonomous driving dataset, (2020). 10.48550\/arXiv.2004.06320."},{"key":"10.1016\/j.inffus.2026.104543_bib0203","series-title":"2019 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"8740","article-title":"Argoverse: 3D tracking and forecasting with rich maps","author":"Chang","year":"2019"},{"key":"10.1016\/j.inffus.2026.104543_bib0204","unstructured":"N. Corporation, PhysicalAI: autonomous vehicles dataset, 2025. Published: Hugging Face, https:\/\/huggingface.co\/datasets\/nvidia\/PhysicalAI-Autonomous-Vehicles."},{"key":"10.1016\/j.inffus.2026.104543_bib0205","series-title":"2020 IEEE International Conference on Robotics and Automation (ICRA)","first-page":"6433","article-title":"The Oxford radar robotcar dataset: a radar extension to the Oxford robotcar dataset","author":"Barnes","year":"2020"},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0206","doi-asserted-by":"crossref","first-page":"3","DOI":"10.1177\/0278364916679498","article-title":"1 Year, 1000\u202fkm: the Oxford robotcar dataset","volume":"36","author":"Maddern","year":"2017","journal-title":"Int. J. Robot. Res."},{"key":"10.1016\/j.inffus.2026.104543_bib0207","doi-asserted-by":"crossref","unstructured":"G. Ros, L. Sellart, J. Materzynska, D. Vazquez, A.M. Lopez, The synthia dataset: a large collection of synthetic images for semantic segmentation of urban scenes, 2016, pp. 3234\u20133243. https:\/\/www.cv-foundation.org\/openaccess\/content_cvpr_2016\/html\/Ros_The_SYNTHIA_Dataset_CVPR_2016_paper.html.","DOI":"10.1109\/CVPR.2016.352"},{"key":"10.1016\/j.inffus.2026.104543_bib0208","doi-asserted-by":"crossref","unstructured":"A. Gaidon, Q. Wang, Y. Cabon, E. Vig, Virtual worlds as proxy for multi-object tracking analysis, 2016, pp. 4340\u20134349. https:\/\/www.cv-foundation.org\/openaccess\/content_cvpr_2016\/html\/Gaidon_Virtual_Worlds_as_CVPR_2016_paper.html.","DOI":"10.1109\/CVPR.2016.470"},{"key":"10.1016\/j.inffus.2026.104543_sbref0209","series-title":"Proceedings of the 1st Annual Conference on Robot Learning","first-page":"1","article-title":"CARLA: an open urban driving simulator","author":"Dosovitskiy","year":"2017"},{"key":"10.1016\/j.inffus.2026.104543_sbref0210","series-title":"World Simulation with Video Foundation Models for Physical AI","author":"Liu","year":"2025"},{"issue":"1","key":"10.1016\/j.inffus.2026.104543_bib0211","doi-asserted-by":"crossref","first-page":"99","DOI":"10.1145\/3503250","article-title":"NeRF: representing scenes as neural radiance fields for view synthesis","volume":"65","author":"Mildenhall","year":"2022","journal-title":"Commun. ACM"},{"issue":"4","key":"10.1016\/j.inffus.2026.104543_bib0212","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1145\/3592433","article-title":"3D Gaussian splatting for real-time radiance field rendering","volume":"42","author":"Kerbl","year":"2023","journal-title":"ACM Trans. Graph."},{"key":"10.1016\/j.inffus.2026.104543_bib0213","doi-asserted-by":"crossref","unstructured":"X. Zhou, Z. Lin, X. Shan, Y. Wang, D. Sun, M.-H. Yang, DrivingGaussian: composite Gaussian splatting for surrounding dynamic autonomous driving scenes, 2024, pp. 21634\u201321643. https:\/\/openaccess.thecvf.com\/content\/CVPR2024\/html\/Zhou_DrivingGaussian_Composite_Gaussian_Splatting_for_Surrounding_Dynamic_Autonomous_Driving_Scenes_CVPR_2024_paper.html.","DOI":"10.1109\/CVPR52733.2024.02044"},{"key":"10.1016\/j.inffus.2026.104543_bib0214","unstructured":"NVIDIA, NVIDIA omniverse NuRec libraries, 2026, https:\/\/docs.nvidia.com\/nurec\/index.html. Accessed 2026-04-16."},{"key":"10.1016\/j.inffus.2026.104543_bib0215","unstructured":"P.D. Inc, Parallel domain- software augmented testing for autonomous systems, 2025. https:\/\/paralleldomain.com\/."},{"key":"10.1016\/j.inffus.2026.104543_bib0216","series-title":"Int. Conf. Comput. Vis.","first-page":"8515","article-title":"Towards viewpoint robustness in bird\u2019s eye view segmentation","author":"Klinghoffer","year":"2023"},{"key":"10.1016\/j.inffus.2026.104543_bib0217","series-title":"Proc. Conf. Robot Learn.","first-page":"4698","article-title":"DriveVLM: the convergence of autonomous driving and large vision-language models","volume":"270","author":"Tian","year":"2025"}],"container-title":["Information Fusion"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S1566253526004215?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S1566253526004215?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2026,6,26]],"date-time":"2026-06-26T14:44:55Z","timestamp":1782485095000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S1566253526004215"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,12]]},"references-count":217,"alternative-id":["S1566253526004215"],"URL":"https:\/\/doi.org\/10.1016\/j.inffus.2026.104543","relation":{},"ISSN":["1566-2535"],"issn-type":[{"value":"1566-2535","type":"print"}],"subject":[],"published":{"date-parts":[[2026,12]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"Toward a universal perception layer: A survey on sensor-agnostic advanced driver assistance systems","name":"articletitle","label":"Article Title"},{"value":"Information Fusion","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1016\/j.inffus.2026.104543","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"\u00a9 2026 Elsevier B.V. All rights are reserved, including those for text and data mining, AI training, and similar technologies.","name":"copyright","label":"Copyright"}],"article-number":"104543"}}