{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,9]],"date-time":"2026-05-09T05:26:57Z","timestamp":1778304417424,"version":"3.51.4"},"reference-count":41,"publisher":"IEEE","license":[{"start":{"date-parts":[[2023,10,1]],"date-time":"2023-10-01T00:00:00Z","timestamp":1696118400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2023,10,1]],"date-time":"2023-10-01T00:00:00Z","timestamp":1696118400000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"}],"content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2023,10,1]]},"DOI":"10.1109\/iros55552.2023.10342434","type":"proceedings-article","created":{"date-parts":[[2023,12,13]],"date-time":"2023-12-13T14:17:55Z","timestamp":1702477075000},"page":"7704-7710","source":"Crossref","is-referenced-by-count":12,"title":["Self-Supervised Event-Based Monocular Depth Estimation Using Cross-Modal Consistency"],"prefix":"10.1109","author":[{"given":"Junyu","family":"Zhu","sequence":"first","affiliation":[{"name":"Institute of Cyber-Systems and Control, Zhejiang University,Hangzhou,China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Lina","family":"Liu","sequence":"additional","affiliation":[{"name":"Institute of Cyber-Systems and Control, Zhejiang University,Hangzhou,China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Bofeng","family":"Jiang","sequence":"additional","affiliation":[{"name":"Institute of Cyber-Systems and Control, Zhejiang University,Hangzhou,China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Feng","family":"Wen","sequence":"additional","affiliation":[{"name":"Huawei Technologies,Noah&#x0027;s Ark Lab,Beijing,China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hongbo","family":"Zhang","sequence":"additional","affiliation":[{"name":"Huawei Technologies,Noah&#x0027;s Ark Lab,Beijing,China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Wanlong","family":"Li","sequence":"additional","affiliation":[{"name":"Huawei Technologies,Noah&#x0027;s Ark Lab,Beijing,China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yong","family":"Liu","sequence":"additional","affiliation":[{"name":"Institute of Cyber-Systems and Control, Zhejiang University,Hangzhou,China"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"263","reference":[{"key":"ref1","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2019.2963386"},{"key":"ref2","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52688.2022.00569"},{"key":"ref3","doi-asserted-by":"publisher","DOI":"10.1109\/3DV53792.2021.00030"},{"key":"ref4","doi-asserted-by":"publisher","DOI":"10.1109\/3DV50981.2020.00063"},{"key":"ref5","doi-asserted-by":"publisher","DOI":"10.1109\/LRA.2021.3060707"},{"key":"ref6","article-title":"Adabins: Depth estimation using adaptive bins","volume-title":"CVPR","author":"Bhat","year":"2020"},{"key":"ref7","doi-asserted-by":"publisher","DOI":"10.1109\/cvpr52688.2022.00389"},{"key":"ref8","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2019.00393"},{"key":"ref9","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.238"},{"key":"ref10","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2019.2936024"},{"key":"ref11","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2018.00407"},{"key":"ref12","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00108"},{"key":"ref13","doi-asserted-by":"publisher","DOI":"10.1109\/IROS45743.2020.9341224"},{"key":"ref14","article-title":"Depth map prediction from a single image using a multi-scale deep network","author":"Eigen","year":"2014","journal-title":"NeurIPS"},{"key":"ref15","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2018.00214"},{"key":"ref16","article-title":"An image is worth 16x16 words: Transformers for image recognition at scale","volume-title":"International Conference for Learning Representations (ICLR)","author":"Dosovitskiy","year":"2021"},{"key":"ref17","author":"Mei","year":"2021","journal-title":"Transvos: Video object segmentation with transformers"},{"key":"ref18","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"ref19","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-319-46484-8_45"},{"key":"ref20","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.699"},{"key":"ref21","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2017.700"},{"key":"ref22","doi-asserted-by":"publisher","DOI":"10.1109\/IROS45743.2020.9340802"},{"key":"ref23","doi-asserted-by":"publisher","DOI":"10.1609\/aaai.v35i3.16329"},{"key":"ref24","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR42600.2020.01313"},{"key":"ref25","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-58565-5_35"},{"key":"ref26","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV48922.2021.01527"},{"key":"ref27","article-title":"Ra-depth: Resolution adaptive self-supervised monocular depth estimation","volume-title":"ECCV","author":"Mu","year":"2022"},{"key":"ref28","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR52688.2022.00163"},{"key":"ref29","doi-asserted-by":"publisher","DOI":"10.1609\/aaai.v33i01.33018001"},{"key":"ref30","article-title":"3d packing for self-supervised monocular depth estimation","author":"P","year":"2020","journal-title":"CVPR"},{"key":"ref31","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-030-01246-5_15"},{"key":"ref32","doi-asserted-by":"publisher","DOI":"10.1109\/TIM.2022.3144229"},{"key":"ref33","doi-asserted-by":"publisher","DOI":"10.15607\/RSS.2018.XIV.062"},{"key":"ref34","doi-asserted-by":"publisher","DOI":"10.1109\/TNNLS.2013.2273537"},{"key":"ref35","first-page":"0","article-title":"Unsupervised event-based optical flow using motion compensation","volume-title":"Proceedings of the European Conference on Computer Vision (ECCV) Workshops","author":"Zihao Zhu","year":"2018"},{"key":"ref36","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2018.00186"},{"key":"ref37","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2019.00573"},{"key":"ref38","doi-asserted-by":"publisher","DOI":"10.1109\/TIP.2003.819861"},{"key":"ref39","article-title":"Self-supervised monocular depth estimation with internal feature fusion","volume-title":"British Machine Vision Conference (BMVC)","author":"Zhou","year":"2021"},{"key":"ref40","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP40776.2020.9053405"},{"key":"ref41","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV48922.2021.00422"}],"event":{"name":"2023 IEEE\/RSJ International Conference on Intelligent Robots and Systems (IROS)","location":"Detroit, MI, USA","start":{"date-parts":[[2023,10,1]]},"end":{"date-parts":[[2023,10,5]]}},"container-title":["2023 IEEE\/RSJ International Conference on Intelligent Robots and Systems (IROS)"],"original-title":[],"link":[{"URL":"http:\/\/xplorestaging.ieee.org\/ielx7\/10341341\/10341342\/10342434.pdf?arnumber=10342434","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,1,11]],"date-time":"2024-01-11T20:32:22Z","timestamp":1705005142000},"score":1,"resource":{"primary":{"URL":"https:\/\/ieeexplore.ieee.org\/document\/10342434\/"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,10,1]]},"references-count":41,"URL":"https:\/\/doi.org\/10.1109\/iros55552.2023.10342434","relation":{},"subject":[],"published":{"date-parts":[[2023,10,1]]}}}