{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,17]],"date-time":"2026-05-17T09:09:28Z","timestamp":1779008968704,"version":"3.51.4"},"publisher-location":"New York, NY, USA","reference-count":90,"publisher":"ACM","license":[{"start":{"date-parts":[[2026,5,10]],"date-time":"2026-05-10T00:00:00Z","timestamp":1778371200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0\/legalcode"}],"funder":[{"name":"Ministry of Education (MOE), Singapore","award":["T2EP20124-0055"],"award-info":[{"award-number":["T2EP20124-0055"]}]},{"name":"National Research Foundation, Singapore","award":["NRF-NRFI05-2019-0007"],"award-info":[{"award-number":["NRF-NRFI05-2019-0007"]}]}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2026,5,11]]},"DOI":"10.1145\/3774906.3802747","type":"proceedings-article","created":{"date-parts":[[2026,5,8]],"date-time":"2026-05-08T14:20:14Z","timestamp":1778250014000},"page":"1002-1015","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":0,"title":["FusionBridge: Enhancing Multi-View Multi-Modal Sensing and Perception for Edge Intelligence"],"prefix":"10.1145","author":[{"ORCID":"https:\/\/orcid.org\/0009-0005-5270-1794","authenticated-orcid":false,"given":"Dhanuja","family":"Wanniarachchige","sequence":"first","affiliation":[{"name":"Singapore Management University, Singapore, Singapore"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-6483-5037","authenticated-orcid":false,"given":"Kasthuri","family":"Jayarajah","sequence":"additional","affiliation":[{"name":"New Jersey Institute of Technology, Newark, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-3883-7220","authenticated-orcid":false,"given":"Tarek","family":"Abdelzaher","sequence":"additional","affiliation":[{"name":"UIUC, Urbana-Champaign, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-1212-1769","authenticated-orcid":false,"given":"Archan","family":"Misra","sequence":"additional","affiliation":[{"name":"Singapore Management University, Singapore, Singapore"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"320","published-online":{"date-parts":[[2026,5,10]]},"reference":[{"key":"e_1_3_3_2_2_2","unstructured":"Samira Abnar and Willem Zuidema. 2020. Quantifying attention flow in transformers. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2005.00928 (2020)."},{"key":"e_1_3_3_2_3_2","unstructured":"Davide Allegro Matteo Terreran and Stefano Ghidoni. 2025. Calib3R: A 3D Foundation Model for Multi-Camera to Robot Calibration and 3D Metric-Scaled Scene Reconstruction. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2509.08813 (2025)."},{"key":"e_1_3_3_2_4_2","doi-asserted-by":"crossref","unstructured":"Nicholas\u00a0I Arnold Plamen Angelov and Peter\u00a0M Atkinson. 2022. An improved explainable point cloud classifier (XPCC). IEEE Transactions on Artificial Intelligence 4 1 (2022) 71\u201380.","DOI":"10.1109\/TAI.2022.3150647"},{"key":"e_1_3_3_2_5_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52688.2022.00116"},{"key":"e_1_3_3_2_6_2","doi-asserted-by":"crossref","unstructured":"Karen Byrd Alei Fan EunSol Her Yiran Liu Barbara Almanza and Stephen Leitch. 2021. Robot vs human: expectations performances and gaps in off-premise restaurant service modes. International Journal of Contemporary Hospitality Management 33 11 (2021) 3996\u20134016.","DOI":"10.1108\/IJCHM-07-2020-0721"},{"key":"e_1_3_3_2_7_2","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"e_1_3_3_2_8_2","unstructured":"T. Chavdarova P. Baqu\u00e9 S. Bouquet A. Maksai C. Jose L. Lettry P. Fua L. Van\u00a0Gool and F. Fleuret. 2017. The WILDTRACK Multi-Camera Person Dataset. CoRR abs\/1707.09299 (2017). http:\/\/arxiv.org\/abs\/1707.09299"},{"key":"e_1_3_3_2_9_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR46437.2021.00084"},{"key":"e_1_3_3_2_10_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.691"},{"key":"e_1_3_3_2_11_2","doi-asserted-by":"crossref","unstructured":"Yingfeng Chen Feng Wu Wei Shuai and Xiaoping Chen. 2017. Robots serve humans in public places\u2014KeJia robot as a shopping assistant. International Journal of Advanced Robotic Systems 14 3 (2017) 1729881417703569.","DOI":"10.1177\/1729881417703569"},{"key":"e_1_3_3_2_12_2","doi-asserted-by":"crossref","unstructured":"Haixing Cheng Chengyong Liu Wenzhe Gu Yuyi Wu Mengye Zhao Wentao Liu and Naibang Wang. 2025. LGMMFusion: A LiDAR-guided multi-modal fusion framework for enhanced 3D object detection. PloS one 20 9 (2025) e0331195.","DOI":"10.1371\/journal.pone.0331195"},{"key":"e_1_3_3_2_13_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2019.00905"},{"key":"e_1_3_3_2_14_2","doi-asserted-by":"publisher","DOI":"10.1109\/ITSC.2019.8917126"},{"key":"e_1_3_3_2_15_2","volume-title":"Proceedings of the 2019 Conference of the North American Chapter of the Association for Computational Linguistics (NAACL-HLT)","author":"Devlin Jacob","year":"2019","unstructured":"Jacob Devlin, Ming-Wei Chang, Kenton Lee, and Kristina Toutanova. 2019. BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding. In Proceedings of the 2019 Conference of the North American Chapter of the Association for Computational Linguistics (NAACL-HLT)."},{"key":"e_1_3_3_2_16_2","unstructured":"Alexey Dosovitskiy Lucas Beyer Alexander Kolesnikov Dirk Weissenborn Xiaohua Zhai Thomas Unterthiner Mostafa Dehghani Matthias Minderer Georg Heigold Sylvain Gelly et\u00a0al. 2020. An image is worth 16x16 words: Transformers for image recognition at scale. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2010.11929 (2020)."},{"key":"e_1_3_3_2_17_2","unstructured":"Rachel\u00a0Lea Draelos and Lawrence Carin. 2021. Use HiResCAM instead of Grad\u2010CAM for faithful explanations of convolutional neural networks. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2011.08891 (2021)."},{"key":"e_1_3_3_2_18_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2019.00667"},{"key":"e_1_3_3_2_19_2","unstructured":"Salehe\u00a0Erfanian Ebadi You-Cyuan Jhang Alex Zook Saurav Dhakad Adam Crespi Pete Parisi Steven Borkman Jonathan Hogins and Sujoy Ganguly. 2021. PeopleSansPeople: a synthetic data generator for human-centric computer vision. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2112.09290 (2021)."},{"key":"e_1_3_3_2_20_2","doi-asserted-by":"crossref","unstructured":"M. Everingham L. Van\u00a0Gool C. Williams J. Winn and A. Zisserman. 2010. The pascal visual object classes (voc) challenge. International journal of computer vision 88 2 (2010) 303\u2013338.","DOI":"10.1007\/s11263-009-0275-4"},{"key":"e_1_3_3_2_21_2","unstructured":"Yuxin Fang Bencheng Liao Xinggang Wang Jiemin Fang Jiyang Qi Rui Wu Jianwei Niu and Wenyu Liu. 2021. You only look at one sequence: Rethinking transformer in vision through object detection. Advances in Neural Information Processing Systems 34 (2021) 26183\u201326197."},{"key":"e_1_3_3_2_22_2","unstructured":"Pierluigi Ferrari. 2023. SSD Keras (VGG16). https:\/\/github.com\/pierluigiferrari\/ssd_keras. Accessed: June 1 2024."},{"key":"e_1_3_3_2_23_2","doi-asserted-by":"publisher","DOI":"10.1109\/PETS-WINTER.2009.5399556"},{"key":"e_1_3_3_2_24_2","doi-asserted-by":"crossref","unstructured":"Juan\u00a0Miguel Garcia-Haro Edwin\u00a0Daniel O\u00f1a Juan Hernandez-Vicen Santiago Martinez and Carlos Balaguer. 2020. Service robots in catering applications: A review and future challenges. Electronics 10 1 (2020) 47.","DOI":"10.3390\/electronics10010047"},{"key":"e_1_3_3_2_25_2","doi-asserted-by":"crossref","unstructured":"Andreas Geiger Philip Lenz Christoph Stiller and Raquel Urtasun. 2013. Vision meets robotics: The kitti dataset. The international journal of robotics research 32 11 (2013) 1231\u20131237.","DOI":"10.1177\/0278364913491297"},{"key":"e_1_3_3_2_26_2","doi-asserted-by":"publisher","DOI":"10.5555\/2354409.2354978"},{"key":"e_1_3_3_2_27_2","doi-asserted-by":"crossref","unstructured":"Lipeng Gu Shaoyuan Sun Xunhua Liu and Xiang Li. 2021. Centertrack3d: Improved centertrack more suitable for three-dimensional objects. Journal of Autonomous Vehicles and Systems 1 2 (2021) 021004.","DOI":"10.1115\/1.4050863"},{"key":"e_1_3_3_2_28_2","doi-asserted-by":"publisher","DOI":"10.5555\/861369"},{"key":"e_1_3_3_2_29_2","doi-asserted-by":"crossref","unstructured":"Shmuel\u00a0Y Hayoun Meir Halachmi Doron Serebro Kfir Twizer Elinor Medezinski Liron Korkidi Moshik Cohen and Itai Orr. 2024. Physics and semantic informed multi-sensor calibration via optimization theory and self-supervised learning. Scientific Reports 14 1 (2024) 2541.","DOI":"10.1038\/s41598-024-53009-z"},{"key":"e_1_3_3_2_30_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2017.322"},{"key":"e_1_3_3_2_31_2","doi-asserted-by":"crossref","unstructured":"Jane Holland Liz Kingston Conor McCarthy Eddie Armstrong Peter O\u2019Dwyer Fionn Merz and Mark McConnell. 2021. Service robots in the healthcare sector. Robotics 10 1 (2021) 47.","DOI":"10.3390\/robotics10010047"},{"key":"e_1_3_3_2_32_2","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58571-6_1"},{"key":"e_1_3_3_2_33_2","volume-title":"Advances in Neural Information Processing Systems (NeurIPS)","author":"Hu Xiaoyu","year":"2022","unstructured":"Xiaoyu Hu, Yiming Liu, Hao Chen, et\u00a0al. 2022. Where2comm: Communication-efficient collaborative perception via spatial confidence maps. In Advances in Neural Information Processing Systems (NeurIPS)."},{"key":"e_1_3_3_2_34_2","unstructured":"Junjie Huang Guan Huang Zheng Zhu Yun Ye and Dalong Du. 2021. Bevdet: High-performance multi-camera 3d object detection in bird-eye-view. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2112.11790 (2021)."},{"key":"e_1_3_3_2_35_2","unstructured":"Keli Huang Botian Shi Xiang Li Xin Li Siyuan Huang and Yikang Li. 2022. Multi-modal sensor fusion for auto driving perception: A survey. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2202.02703 (2022)."},{"key":"e_1_3_3_2_36_2","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58555-6_3"},{"key":"e_1_3_3_2_37_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52734.2025.02551"},{"key":"e_1_3_3_2_38_2","volume-title":"Proceedings of the 2022 IEEE International Conference on Computer Communications (INFOCOM), Virtual Conference, May","author":"JAYARAJAH Kasthuri","unstructured":"Kasthuri JAYARAJAH, Dhanuja WANNIARACHCHIGE, Tarek ABDELZAHER, and Archan MISRA. [n. d.]. ComAI: Enabling lightweight, collaborative intelligence by retrofitting vision DNNs.(2022). In Proceedings of the 2022 IEEE International Conference on Computer Communications (INFOCOM), Virtual Conference, May."},{"key":"e_1_3_3_2_39_2","volume-title":"Ultralytics YOLOv8","author":"Jocher Glenn","year":"2023","unstructured":"Glenn Jocher, Ayush Chaurasia, and Jing Qiu. 2023. Ultralytics YOLOv8. https:\/\/github.com\/ultralytics\/ultralytics"},{"key":"e_1_3_3_2_40_2","unstructured":"Alex Krizhevsky Ilya Sutskever and Geoffrey\u00a0E Hinton. 2012. Imagenet classification with deep convolutional neural networks. Advances in neural information processing systems 25 (2012)."},{"key":"e_1_3_3_2_41_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.01298"},{"key":"e_1_3_3_2_42_2","unstructured":"Laura Leal-Taix\u00e9 Anton Milan Ian Reid Stefan Roth and Konrad Schindler. 2015. Motchallenge 2015: Towards a benchmark for multi-target tracking. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/1504.01942 (2015)."},{"key":"e_1_3_3_2_43_2","doi-asserted-by":"crossref","unstructured":"Youn\u00a0Joo Lee Jun\u00a0Young Hwang Jiwon Park Ho\u00a0Gi Jung and Jae\u00a0Kyu Suhr. 2024. Deep neural network-based flood monitoring system fusing rgb and lwir cameras for embedded iot edge devices. Remote Sensing 16 13 (2024) 2358.","DOI":"10.3390\/rs16132358"},{"key":"e_1_3_3_2_44_2","unstructured":"Yixing Li and Fengbo Ren. 2019. Light-weight retinanet for object detection. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/1905.10011 (2019)."},{"key":"e_1_3_3_2_45_2","unstructured":"Zhiqi Li Wenhai Wang Hongyang Li Enze Xie Chonghao Sima Tong Lu Qiao Yu and Jifeng Dai. 2024. Bevformer: learning bird\u2019s-eye-view representation from lidar-camera via spatiotemporal transformers. IEEE Transactions on Pattern Analysis and Machine Intelligence (2024)."},{"key":"e_1_3_3_2_46_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2017.324"},{"key":"e_1_3_3_2_47_2","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-319-46448-0_2"},{"key":"e_1_3_3_2_48_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00062"},{"key":"e_1_3_3_2_49_2","doi-asserted-by":"publisher","DOI":"10.1145\/3123266.3123436"},{"key":"e_1_3_3_2_50_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV48922.2021.00363"},{"key":"e_1_3_3_2_51_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV48922.2021.00290"},{"key":"e_1_3_3_2_52_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICAL.2009.5262912"},{"key":"e_1_3_3_2_53_2","doi-asserted-by":"publisher","DOI":"10.1109\/IROS45743.2020.9341791"},{"key":"e_1_3_3_2_54_2","unstructured":"Charles\u00a0Ruizhongtai Qi Li Yi Hao Su and Leonidas\u00a0J Guibas. 2017. Pointnet++: Deep hierarchical feature learning on point sets in a metric space. Advances in neural information processing systems 30 (2017)."},{"key":"e_1_3_3_2_55_2","first-page":"389","volume-title":"2024 International Conference on Digital Image Computing: Techniques and Applications (DICTA)","author":"Qiao Donghao","year":"2024","unstructured":"Donghao Qiao, Farhana Zulkernine, and Aman Anand. 2024. Cobevfusion cooperative perception with lidar-camera bird\u2019s eye view fusion. In 2024 International Conference on Digital Image Computing: Techniques and Applications (DICTA). IEEE, 389\u2013396."},{"key":"e_1_3_3_2_56_2","first-page":"8748","volume-title":"International conference on machine learning","author":"Radford Alec","year":"2021","unstructured":"Alec Radford, Jong\u00a0Wook Kim, Chris Hallacy, Aditya Ramesh, Gabriel Goh, Sandhini Agarwal, Girish Sastry, Amanda Askell, Pamela Mishkin, Jack Clark, et\u00a0al. 2021. Learning transferable visual models from natural language supervision. In International conference on machine learning. PmLR, 8748\u20138763."},{"key":"e_1_3_3_2_57_2","doi-asserted-by":"crossref","unstructured":"Darshana Rathnayake Meera Radhakrishnan Inseok Hwang and Archan Misra. 2024. LILOC: Leveraging LiDARs for accurate 3D localization in dynamic indoor environments. ACM Transactions on Internet of Things 5 4 (2024) 1\u201333.","DOI":"10.1145\/3695881"},{"key":"e_1_3_3_2_58_2","unstructured":"Shaoqing Ren Kaiming He Ross Girshick and Jian Sun. 2015. Faster r-cnn: Towards real-time object detection with region proposal networks. Advances in neural information processing systems 28 (2015)."},{"key":"e_1_3_3_2_59_2","doi-asserted-by":"crossref","unstructured":"Shaoqing Ren Kaiming He Ross Girshick and Jian Sun. 2016. Faster R-CNN: Towards real-time object detection with region proposal networks. IEEE transactions on pattern analysis and machine intelligence 39 6 (2016) 1137\u20131149.","DOI":"10.1109\/TPAMI.2016.2577031"},{"key":"e_1_3_3_2_60_2","unstructured":"Tianhe Ren Qing Jiang Shilong Liu Zhaoyang Zeng Wenlong Liu Han Gao Hongjie Huang Zhengyu Ma Xiaoke Jiang Yihao Chen et\u00a0al. 2024. Grounding dino 1.5: Advance the\" edge\" of open-set object detection. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2405.10300 (2024)."},{"key":"e_1_3_3_2_61_2","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-319-48881-3_2"},{"key":"e_1_3_3_2_62_2","unstructured":"Clearpath Robotics. 2013. TurtleBot 2. https:\/\/www.turtlebot.com\/turtlebot2\/. Accessed: June 1 2024."},{"key":"e_1_3_3_2_63_2","doi-asserted-by":"publisher","DOI":"10.1145\/1089444.1089474"},{"key":"e_1_3_3_2_64_2","doi-asserted-by":"publisher","DOI":"10.21236\/ADA164453"},{"key":"e_1_3_3_2_65_2","doi-asserted-by":"crossref","unstructured":"Yongxin Shao Zhetao Sun Aihong Tan and Tianhong Yan. 2023. Efficient three-dimensional point cloud object detection based on improved Complex-YOLO. Frontiers in Neurorobotics 17 (2023) 1092564.","DOI":"10.3389\/fnbot.2023.1092564"},{"key":"e_1_3_3_2_66_2","doi-asserted-by":"crossref","unstructured":"Sachin Sharma Richard\u00a0T Meyer and Zachary\u00a0D Asher. 2024. AEPF: Attention-Enabled Point Fusion for 3D Object Detection. Sensors 24 17 (2024) 5841.","DOI":"10.3390\/s24175841"},{"key":"e_1_3_3_2_67_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR42600.2020.01054"},{"key":"e_1_3_3_2_68_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00086"},{"key":"e_1_3_3_2_69_2","unstructured":"Martin Simon Stefan Milz Karl Amende and Horst\u2010Michael Gross. 2018. Complex\u2010YOLO: Real\u2010time 3D Object Detection on Point Clouds. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/1803.06199 (2018)."},{"key":"e_1_3_3_2_70_2","first-page":"0","volume-title":"Proceedings of the European conference on computer vision (ECCV) workshops","author":"Simony Martin","year":"2018","unstructured":"Martin Simony, Stefan Milzy, Karl Amendey, and Horst-Michael Gross. 2018. Complex-yolo: An euler-region-proposal for real-time 3d object detection on point clouds. In Proceedings of the European conference on computer vision (ECCV) workshops. 0\u20130."},{"key":"e_1_3_3_2_71_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICRA.2019.8794195"},{"key":"e_1_3_3_2_72_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2016.94"},{"key":"e_1_3_3_2_73_2","doi-asserted-by":"crossref","unstructured":"Simon Speth Artur Goncalves Bastien Rigault Satoshi Suzuki Mondher Bouazizi Yutaka Matsuo and Helmut Prendinger. 2022. Deep learning with RGB and thermal images onboard a drone for monitoring operations. Journal of Field Robotics 39 6 (2022) 840\u2013868.","DOI":"10.1002\/rob.22082"},{"key":"e_1_3_3_2_74_2","unstructured":"Synty Studios. 2023. Polygon City Pack: Environment and Interior. https:\/\/assetstore.unity.com\/packages\/3d\/polygon-city-pack-environment-and-interior-free-101685. Accessed: October 1 2023."},{"key":"e_1_3_3_2_75_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR42600.2020.00252"},{"key":"e_1_3_3_2_76_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR42600.2020.01079"},{"key":"e_1_3_3_2_77_2","doi-asserted-by":"publisher","DOI":"10.1145\/3669721.3669748"},{"key":"e_1_3_3_2_78_2","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2019.00651"},{"key":"e_1_3_3_2_79_2","unstructured":"Zhi Tian Chunhua Shen Hao Chen and Tong He. 2020. FCOS: A simple and strong anchor-free object detector. IEEE transactions on pattern analysis and machine intelligence 44 4 (2020) 1922\u20131933."},{"key":"e_1_3_3_2_80_2","unstructured":"Ashish Vaswani Noam Shazeer Niki Parmar Jakob Uszkoreit Llion Jones Aidan\u00a0N Gomez \u0141ukasz Kaiser and Illia Polosukhin. 2017. Attention is all you need. Advances in neural information processing systems 30 (2017)."},{"key":"e_1_3_3_2_81_2","unstructured":"Velodyne LiDAR. 2016. Velodyne Puck VLP-16. https:\/\/velodynelidar.com\/products\/puck\/."},{"key":"e_1_3_3_2_82_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR42600.2020.00466"},{"key":"e_1_3_3_2_83_2","doi-asserted-by":"publisher","DOI":"10.1109\/IVS.2018.8500387"},{"key":"e_1_3_3_2_84_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR46437.2021.01293"},{"key":"e_1_3_3_2_85_2","unstructured":"Zhenbo Xu Wei Zhang Xiao Tan Wei Yang Xiangbo Su Yuchen Yuan Hongwu Zhang Shilei Wen Errui Ding and Liusheng Huang. 2020. Pointtrack++ for effective online multi-object tracking and segmentation. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2007.01549 (2020)."},{"key":"e_1_3_3_2_86_2","doi-asserted-by":"crossref","unstructured":"Yan Yan Yuxing Mao and Bo Li. 2018. Second: Sparsely embedded convolutional detection. Sensors 18 10 (2018) 3337.","DOI":"10.3390\/s18103337"},{"key":"e_1_3_3_2_87_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR46437.2021.01161"},{"key":"e_1_3_3_2_88_2","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58583-9_43"},{"key":"e_1_3_3_2_89_2","unstructured":"Binyu Zhao Wei Zhang and Zhaonian Zou. 2023. Bm2cp: Efficient collaborative perception with lidar-camera modalities. arXiv preprint arXiv:https:\/\/arXiv.org\/abs\/2310.14702 (2023)."},{"key":"e_1_3_3_2_90_2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2018.00472"},{"key":"e_1_3_3_2_91_2","doi-asserted-by":"crossref","unstructured":"Zijie Zhou Jingyi Xu Guangming Xiong and Junyi Ma. 2023. Lcpr: A multi-scale attention-based lidar-camera fusion network for place recognition. IEEE Robotics and Automation Letters 9 2 (2023) 1342\u20131349.","DOI":"10.1109\/LRA.2023.3346753"}],"event":{"name":"SenSys '26: ACM\/IEEE International Conference on Embedded Artificial Intelligence and Sensing Systems","location":"Saint Malo France","acronym":"SenSys '26","sponsor":["SIGBED ACM Special Interest Group on Embedded Systems","SIGMOBILE ACM Special Interest Group on Mobility of Systems, Users, Data and Computing","IEEE CS"]},"container-title":["Proceedings of the 2026 ACM\/IEEE International Conference on Embedded Artificial Intelligence and Sensing Systems"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3774906.3802747","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,5,17]],"date-time":"2026-05-17T08:30:17Z","timestamp":1779006617000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3774906.3802747"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,10]]},"references-count":90,"alternative-id":["10.1145\/3774906.3802747","10.1145\/3774906"],"URL":"https:\/\/doi.org\/10.1145\/3774906.3802747","relation":{},"subject":[],"published":{"date-parts":[[2026,5,10]]},"assertion":[{"value":"2026-05-10","order":3,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}