{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,26]],"date-time":"2026-05-26T11:05:45Z","timestamp":1779793545646,"version":"3.53.1"},"reference-count":35,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T00:00:00Z","timestamp":1775001600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T00:00:00Z","timestamp":1775001600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Special Deep Space Exploration Cultivation Fund of USTC","award":["YD2390000601"],"award-info":[{"award-number":["YD2390000601"]}]},{"name":"University-Enterprise Joint Scientific Research and Talent Training Project of the Advanced Research Institute of Anhui Provinc"},{"name":"Fund of Robot Technology Used for Special Environment Key Laboratory of Sichuan Province","award":["25kftk03"],"award-info":[{"award-number":["25kftk03"]}]},{"name":"USTC-JAC Smart Electric Vehicle Joint Lab"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2026,4]]},"DOI":"10.1007\/s10489-026-07242-w","type":"journal-article","created":{"date-parts":[[2026,4,24]],"date-time":"2026-04-24T12:41:14Z","timestamp":1777034474000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["DGKD: Depth-Guided knowledge distillation network for monocular 3D object detection"],"prefix":"10.1007","volume":"56","author":[{"given":"Xinyu","family":"Zhang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5688-4130","authenticated-orcid":false,"given":"Qiang","family":"Ling","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,4,24]]},"reference":[{"key":"7242_CR1","doi-asserted-by":"crossref","unstructured":"Caesar H, Bankiti V, Lang AH, Vora S, Liong VE, Xu Q, Krishnan A, Pan Y, Baldan G, Beijbom O (2019) nuscenes: A multimodal dataset for autonomous driving. In: 2020 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 11618\u201311628","DOI":"10.1109\/CVPR42600.2020.01164"},{"key":"7242_CR2","doi-asserted-by":"crossref","unstructured":"Ettinger SM, Cheng S, Caine B, Liu C, Zhao H, Pradhan S, Chai Y, Sapp B, Qi C, Zhou Y, Yang Z, Chouard A, Sun P, Ngiam J, Vasudevan V, McCauley A, Shlens J, Anguelov D (2021) Large scale interactive motion forecasting for autonomous driving: The waymo open motion dataset. In: 2021 IEEE\/CVF international conference on computer vision (ICCV), pp 9690\u20139699","DOI":"10.1109\/ICCV48922.2021.00957"},{"issue":"2","key":"7242_CR3","doi-asserted-by":"publisher","first-page":"157","DOI":"10.1177\/0278364907087172","volume":"27","author":"A Saxena","year":"2008","unstructured":"Saxena A, Driemeyer J, Ng AY (2008) Robotic grasping of novel objects using vision. Int J Robot Res 27(2):157\u2013173","journal-title":"Int J Robot Res"},{"issue":"3","key":"7242_CR4","first-page":"1","volume":"55","author":"Y Pan","year":"2024","unstructured":"Pan Y, Chen C, Li D, Zhao Z (2024) A robot path tracking method based on manual guidance and path reinforcement learning: A robot path tracking method based on manual guidance and path reinforcement learning. Appl Intell 55(3):1","journal-title":"Appl Intell"},{"key":"7242_CR5","doi-asserted-by":"crossref","unstructured":"Chen X, Ma H, Wan J, Li B, Xia T (2017) Multi-view 3d object detection network for autonomous driving. In: 2017 IEEE conference on computer vision and pattern recognition (CVPR), pp 6526\u20136534","DOI":"10.1109\/CVPR.2017.691"},{"key":"7242_CR6","doi-asserted-by":"crossref","unstructured":"Shi S, Guo C, Jiang L, Wang Z, Shi J, Wang X, Li H (2019) Pv-rcnn: Point-voxel feature set abstraction for 3d object detection. In: 2020 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 10526\u201310535","DOI":"10.1109\/CVPR42600.2020.01054"},{"issue":"2","key":"7242_CR7","doi-asserted-by":"publisher","first-page":"2362","DOI":"10.1007\/s10489-022-03576-3","volume":"53","author":"F Xu","year":"2022","unstructured":"Xu F, Wang Z, Wang H, Lin L, Liang H (2022) Dynamic vehicle pose estimation and tracking based on motion feedback for lidars. Appl Intell 53(2):2362\u20132390","journal-title":"Appl Intell"},{"key":"7242_CR8","doi-asserted-by":"publisher","first-page":"1259","DOI":"10.1109\/TPAMI.2017.2706685","volume":"40","author":"X Chen","year":"2016","unstructured":"Chen X, Kundu K, Zhu Y, Ma H, Fidler S, Urtasun R (2016) 3d object proposals using stereo imagery for accurate object class detection. IEEE Trans Pattern Anal Mach Intell 40:1259\u20131272","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"7242_CR9","doi-asserted-by":"crossref","unstructured":"K\u00f6nigshof H, Salscheider NO, Stiller C (2019) Realtime 3d object detection for automated driving using stereo vision and semantic information. In: 2019 IEEE intelligent transportation systems conference (ITSC), pp 1405\u20131410","DOI":"10.1109\/ITSC.2019.8917330"},{"key":"7242_CR10","doi-asserted-by":"crossref","unstructured":"Zhang R, Qiu H, Wang T, Guo Z, Tang Y, Cui Z, Xu X, Qiao YJ, Gao P, Li H (2022) Monodetr: Depth-guided transformer for monocular 3d object detection. In: 2023 IEEE\/CVF international conference on computer vision (ICCV), pp 9121\u20139132","DOI":"10.1109\/ICCV51070.2023.00840"},{"key":"7242_CR11","unstructured":"Hinton GE, Vinyals O, Dean J (2015) Distilling the knowledge in a neural network. ArXiv. abs\/1503.02531"},{"issue":"2","key":"7242_CR12","doi-asserted-by":"publisher","first-page":"1997","DOI":"10.1007\/s10489-022-03486-4","volume":"53","author":"C Xu","year":"2022","unstructured":"Xu C, Gao W, Li T, Bai N, Li G, Zhang Y (2022) Teacher-student collaborative knowledge distillation for image classification. Appl Intell 53(2):1997\u20132009","journal-title":"Appl Intell"},{"key":"7242_CR13","unstructured":"Chong Z, Ma X, Zhang H, Yue Y, Li H, Wang Z, Ouyang W (2022) Monodistill: Learning spatial features for monocular 3d object detection. ArXiv. abs\/2201.10830"},{"key":"7242_CR14","doi-asserted-by":"crossref","unstructured":"Mousavian A, Anguelov D, Flynn J, Kosecka J (2016) 3d bounding box estimation using deep learning and geometry. In: 2017 IEEE conference on computer vision and pattern recognition (CVPR), pp 5632\u20135640","DOI":"10.1109\/CVPR.2017.597"},{"key":"7242_CR15","doi-asserted-by":"crossref","unstructured":"Brazil G, Liu X (2019) M3d-rpn: Monocular 3d region proposal network for object detection. In: 2019 IEEE\/CVF international conference on computer vision (ICCV), pp 9286\u20139295","DOI":"10.1109\/ICCV.2019.00938"},{"key":"7242_CR16","unstructured":"Zhou X, Wang D, Kr\u00e4henb\u00fchl P (2019) Objects as points. ArXiv. abs\/1904.07850"},{"key":"7242_CR17","doi-asserted-by":"crossref","unstructured":"Liu Z, Wu Z, T\u2019oth R (2020) Smoke: Single-stage monocular 3d object detection via keypoint estimation. 2020 IEEE\/CVF conference on computer vision and pattern recognition workshops (CVPRW), pp 4289\u20134298","DOI":"10.1109\/CVPRW50498.2020.00506"},{"key":"7242_CR18","doi-asserted-by":"crossref","unstructured":"Zhang Y, Lu J, Zhou J (2021) Objects are different: Flexible monocular 3d object detection. 2021 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 3288\u20133297","DOI":"10.1109\/CVPR46437.2021.00330"},{"key":"7242_CR19","unstructured":"Vaswani A, Shazeer NM, Parmar N, Uszkoreit J, Jones L, Gomez AN, Kaiser L, Polosukhin I (2017) Attention is all you need. In: Neural information processing systems"},{"key":"7242_CR20","doi-asserted-by":"crossref","unstructured":"Huang K-C, Wu T-H, Su H-T, Hsu WH (2022) Monodtr: Monocular 3d object detection with depth-aware transformer. In: 2022 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 4002\u20134011","DOI":"10.1109\/CVPR52688.2022.00398"},{"key":"7242_CR21","doi-asserted-by":"crossref","unstructured":"Zhou Y, Zhu H, Liu Q, Chang S, Guo M (2023) Monoatt: Online monocular 3d object detection with adaptive token transformer. In: 2023 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 17493\u201317503","DOI":"10.1109\/CVPR52729.2023.01678"},{"key":"7242_CR22","unstructured":"Chen G, Choi W, Yu X, Han T, Chandraker M (2017) Learning efficient object detection models with knowledge distillation. In: Proceedings of the 31st international conference on neural information processing systems. NIPS\u201917, pp 742\u2013751. Curran Associates Inc., Red Hook, NY, USA"},{"key":"7242_CR23","doi-asserted-by":"crossref","unstructured":"Guo J, Han K, Wang Y, Wu H, Chen X, Xu C, Xu C (2021) Distilling object detectors via decoupled features. In: 2021 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 2154\u20132164","DOI":"10.1109\/CVPR46437.2021.00219"},{"key":"7242_CR24","doi-asserted-by":"crossref","unstructured":"Guo X, Shi S, Wang X, Li H (2021) Liga-stereo: Learning lidar geometry aware representations for stereo-based 3d detector. In: 2021 IEEE\/CVF international conference on computer vision (ICCV), pp 3133\u20133143","DOI":"10.1109\/ICCV48922.2021.00314"},{"key":"7242_CR25","doi-asserted-by":"publisher","first-page":"87","DOI":"10.1007\/978-3-031-20080-9_6","volume-title":"Computer Vision - ECCV 2022","author":"Y Hong","year":"2022","unstructured":"Hong Y, Dai H, Ding Y (2022) Cross-modality knowledge distillation network for monocular 3d object detection. In: Avidan S, Brostow G, Ciss\u00e9 M, Farinella GM, Hassner T (eds) Computer Vision - ECCV 2022. Springer, Cham, pp 87\u2013104"},{"key":"7242_CR26","doi-asserted-by":"crossref","unstructured":"Geiger A, Lenz P, Urtasun R (2012) Are we ready for autonomous driving? The kitti vision benchmark suite. In: 2012 IEEE conference on computer vision and pattern recognition, pp 3354\u20133361","DOI":"10.1109\/CVPR.2012.6248074"},{"key":"7242_CR27","doi-asserted-by":"crossref","unstructured":"Reading C, Harakeh A, Chae J, Waslander SL (2021) Categorical depth distribution network for monocular 3d object detection. In: 2021 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 8551\u20138560","DOI":"10.1109\/CVPR46437.2021.00845"},{"key":"7242_CR28","doi-asserted-by":"crossref","unstructured":"Ma X, Zhang Y, Xu D, Zhou D, Yi S, Li H, Ouyang W (2021) Delving into localization errors for monocular 3d object detection. In: 2021 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 4719\u20134728","DOI":"10.1109\/CVPR46437.2021.00469"},{"key":"7242_CR29","doi-asserted-by":"crossref","unstructured":"Shi X, Ye Q, Chen X, Chen C, Chen Z, Kim T-K (2021) Geometry-based distance decomposition for monocular 3d object detection. 2021 IEEE\/CVF international conference on computer vision (ICCV), pp 15152\u201315161","DOI":"10.1109\/ICCV48922.2021.01489"},{"key":"7242_CR30","doi-asserted-by":"crossref","unstructured":"Kumar A, Brazil G, Corona E, Parchami A, Liu X (2022) Deviant: Depth equivariant network for monocular 3d object detection. ArXiv. abs\/2207.10758","DOI":"10.1007\/978-3-031-20077-9_39"},{"key":"7242_CR31","doi-asserted-by":"crossref","unstructured":"Peng L, Wu X, Yang Z, Liu H, Cai D (2022) Did-m3d: Decoupling instance depth for monocular 3d object detection. In: European conference on computer vision","DOI":"10.1007\/978-3-031-19769-7_5"},{"key":"7242_CR32","doi-asserted-by":"crossref","unstructured":"Yan L, Yan P, Xiong S, Xiang X, Tan Y (2024) Monocd: Monocular 3d object detection with complementary depths. 2024 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 10248\u201310257","DOI":"10.1109\/CVPR52733.2024.00976"},{"key":"7242_CR33","doi-asserted-by":"crossref","unstructured":"Wu Z, Gan Y, Wu Y, Wang R, Wang X, Pu J (2024) Fd3d: Exploiting foreground depth map for feature-supervised monocular 3d object detection. In: AAAI conference on artificial intelligence","DOI":"10.1609\/aaai.v38i6.28436"},{"key":"7242_CR34","doi-asserted-by":"crossref","unstructured":"Peng L, Xu J, Cheng H, Yang Z, Wu X, Qian W, Wang W, Wu B, Cai D (2023) Learning occupancy for monocular 3d object detection. In: 2024 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp 10281\u201310292","DOI":"10.1109\/CVPR52733.2024.00979"},{"key":"7242_CR35","doi-asserted-by":"crossref","unstructured":"Lu Y, Ma X, Yang L, Zhang T, Liu Y, Chu Q, Yan J, Ouyang W (2021) Geometry uncertainty projection network for monocular 3d object detection. In: 2021 IEEE\/CVF international conference on computer vision (ICCV), pp 3091\u20133101","DOI":"10.1109\/ICCV48922.2021.00310"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-026-07242-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-026-07242-w","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-026-07242-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,5,26]],"date-time":"2026-05-26T10:50:57Z","timestamp":1779792657000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-026-07242-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,4]]},"references-count":35,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2026,4]]}},"alternative-id":["7242"],"URL":"https:\/\/doi.org\/10.1007\/s10489-026-07242-w","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,4]]},"assertion":[{"value":"21 December 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 April 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 April 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing Interests"}},{"value":"Not applicable.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical and Informed Consent for Data Used"}}],"article-number":"227"}}