{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,4]],"date-time":"2026-08-04T00:48:59Z","timestamp":1785804539074,"version":"3.56.0"},"reference-count":71,"publisher":"Springer Science and Business Media LLC","issue":"2","license":[{"start":{"date-parts":[[2025,2,24]],"date-time":"2025-02-24T00:00:00Z","timestamp":1740355200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,2,24]],"date-time":"2025-02-24T00:00:00Z","timestamp":1740355200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"Chongqing Engineering Laboratory for Transportation Engineering Application Robot Open Fund","award":["CELTEAR-KFKT-202003"],"award-info":[{"award-number":["CELTEAR-KFKT-202003"]}]},{"name":"Chongqing Engineering Laboratory for Transportation Engineering Application Robot Open Fund","award":["CELTEAR-KFKT-202003"],"award-info":[{"award-number":["CELTEAR-KFKT-202003"]}]},{"name":"Chongqing Engineering Laboratory for Transportation Engineering Application Robot Open Fund","award":["CELTEAR-KFKT-202003"],"award-info":[{"award-number":["CELTEAR-KFKT-202003"]}]},{"name":"Chongqing Key Laboratory of Urban Rail Transit System Integration and Control Open Fund","award":["CKLURTSIC-KFKT-202006"],"award-info":[{"award-number":["CKLURTSIC-KFKT-202006"]}]},{"name":"Chongqing Key Laboratory of Urban Rail Transit System Integration and Control Open Fund","award":["CKLURTSIC-KFKT-202006"],"award-info":[{"award-number":["CKLURTSIC-KFKT-202006"]}]},{"name":"Chongqing Key Laboratory of Urban Rail Transit System Integration and Control Open Fund","award":["CKLURTSIC-KFKT-202006"],"award-info":[{"award-number":["CKLURTSIC-KFKT-202006"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2025,4]]},"DOI":"10.1007\/s00530-025-01717-5","type":"journal-article","created":{"date-parts":[[2025,2,24]],"date-time":"2025-02-24T08:07:12Z","timestamp":1740384432000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":7,"title":["EATNet: edge-aware and transformer-based network for RGB-D salient object detection"],"prefix":"10.1007","volume":"31","author":[{"given":"Xu","family":"Liu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Chenhua","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xianye","family":"Zhou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Guodong","family":"Fan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,2,24]]},"reference":[{"key":"1717_CR1","doi-asserted-by":"crossref","unstructured":"Wei, Y., Feng, J., Liang, X., Cheng, M.-M., Zhao, Y., Yan, S.: Object region mining with adversarial erasing: A simple classification to semantic segmentation approach. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1568\u20131576 (2017)","DOI":"10.1109\/CVPR.2017.687"},{"key":"1717_CR2","doi-asserted-by":"crossref","unstructured":"Wang, X., You, S., Li, X., Ma, H.: Weakly-supervised semantic segmentation by iteratively mining common object features. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1354\u20131362 (2018)","DOI":"10.1109\/CVPR.2018.00147"},{"key":"1717_CR3","doi-asserted-by":"crossref","unstructured":"Fang, H., Gupta, S., Iandola, F., Srivastava, R.K., Deng, L., Doll\u00e1r, P., Gao, J., He, X., Mitchell, M., Platt, J.C., etal.: From captions to visual concepts and back. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1473\u20131482 (2015)","DOI":"10.1109\/CVPR.2015.7298754"},{"key":"1717_CR4","doi-asserted-by":"publisher","first-page":"90","DOI":"10.1016\/j.cviu.2017.10.001","volume":"163","author":"A Das","year":"2017","unstructured":"Das, A., Agrawal, H., Zitnick, L., Parikh, D., Batra, D.: Human attention in visual question answering: Do humans and deep networks look at the same regions? Comput. Vis. Image Underst. 163, 90\u2013100 (2017)","journal-title":"Comput. Vis. Image Underst."},{"issue":"2","key":"1717_CR5","first-page":"151","volume":"2","author":"S Bi","year":"2014","unstructured":"Bi, S., Li, G., Yu, Y.: Person re-identification using multiple experts with random subspaces. J. Image Graph. 2(2), 151\u2013157 (2014)","journal-title":"J. Image Graph."},{"key":"1717_CR6","doi-asserted-by":"crossref","unstructured":"Zhao, R., Ouyang, W., Wang, X.: Unsupervised salience learning for person re-identification. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3586\u20133593 (2013)","DOI":"10.1109\/CVPR.2013.460"},{"issue":"8","key":"1717_CR7","doi-asserted-by":"publisher","first-page":"2014","DOI":"10.1109\/TVCG.2016.2600594","volume":"23","author":"W Wang","year":"2016","unstructured":"Wang, W., Shen, J., Yu, Y., Ma, K.-L.: Stereoscopic thumbnail creation via efficient stereo saliency detection. IEEE Trans. Vis. Comput. Graph. 23(8), 2014\u20132027 (2016)","journal-title":"IEEE Trans. Vis. Comput. Graph."},{"issue":"7","key":"1717_CR8","doi-asserted-by":"publisher","first-page":"1531","DOI":"10.1109\/TPAMI.2018.2840724","volume":"41","author":"W Wang","year":"2018","unstructured":"Wang, W., Shen, J., Ling, H.: A deep network solution for attention and aesthetics aware photo cropping. IEEE Trans. Pattern Anal. Mach. Intell. 41(7), 1531\u20131544 (2018)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"1717_CR9","doi-asserted-by":"crossref","unstructured":"Wang, W., Zhao, S., Shen, J., Hoi, S.C., Borji, A.: Salient object detection with pyramid attention and salient edges. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1448\u20131457 (2019)","DOI":"10.1109\/CVPR.2019.00154"},{"key":"1717_CR10","doi-asserted-by":"crossref","unstructured":"Wang, W., Shen, J., Cheng, M.-M., Shao, L.: An iterative and cooperative top-down and bottom-up inference network for salient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 5968\u20135977 (2019)","DOI":"10.1109\/CVPR.2019.00612"},{"key":"1717_CR11","doi-asserted-by":"crossref","unstructured":"Liu, J.-J., Hou, Q., Cheng, M.-M., Feng, J., Jiang, J.: A simple pooling-based design for real-time salient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 3917\u20133926 (2019)","DOI":"10.1109\/CVPR.2019.00404"},{"key":"1717_CR12","doi-asserted-by":"publisher","first-page":"376","DOI":"10.1016\/j.patcog.2018.08.007","volume":"86","author":"H Chen","year":"2019","unstructured":"Chen, H., Li, Y., Su, D.: Multi-modal fusion network with multi-scale multi-path and cross-modal interactions for rgb-d salient object detection. Pattern Recogn. 86, 376\u2013385 (2019)","journal-title":"Pattern Recogn."},{"key":"1717_CR13","doi-asserted-by":"crossref","unstructured":"Li, G., Liu, Z., Ye, L., Wang, Y., Ling, H.: Cross-modal weighting network for rgb-d salient object detection. In: European Conference on Computer Vision, pp. 665\u2013681 (2020). Springer","DOI":"10.1007\/978-3-030-58520-4_39"},{"key":"1717_CR14","doi-asserted-by":"crossref","unstructured":"Ji, W., Li, J., Yu, S., Zhang, M., Piao, Y., Yao, S., Bi, Q., Ma, K., Zheng, Y., Lu, H., etal.: Calibrated rgb-d salient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9471\u20139481 (2021)","DOI":"10.1109\/CVPR46437.2021.00935"},{"key":"1717_CR15","doi-asserted-by":"crossref","unstructured":"Wang, K., Liu, C., Wei, H., Jing, L., Zhang, R.: Rfnet: Refined fusion three-branch rgb-d salient object detection network. In: 2024 IEEE International Conference on Image Processing (ICIP), pp. 741\u2013746 (2024). IEEE","DOI":"10.1109\/ICIP51287.2024.10647308"},{"key":"1717_CR16","doi-asserted-by":"publisher","unstructured":"Pang, Y., Zhao, X., Zhang, L., Lu, H.: Caver: Cross-modal view-mixed transformer for bi-modal salient object detection. IEEE Transactions on Image Processing, 1\u20131 (2023). https:\/\/doi.org\/10.1109\/TIP.2023.3234702","DOI":"10.1109\/TIP.2023.3234702"},{"key":"1717_CR17","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., Uszkoreit, J., Houlsby, N.: An image is worth 16x16 words: Transformers for image recognition at scale. ICLR (2021)"},{"key":"1717_CR18","doi-asserted-by":"crossref","unstructured":"Carion, N., Massa, F., Synnaeve, G., Usunier, N., Kirillov, A., Zagoruyko, S.: End-to-end object detection with transformers. In: European Conference on Computer Vision, pp. 213\u2013229 (2020). Springer","DOI":"10.1007\/978-3-030-58452-8_13"},{"key":"1717_CR19","doi-asserted-by":"crossref","unstructured":"Liu, Z., Lin, Y., Cao, Y., Hu, H., Wei, Y., Zhang, Z., Lin, S., Guo, B.: Swin transformer: Hierarchical vision transformer using shifted windows. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV) (2021)","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"1717_CR20","doi-asserted-by":"crossref","unstructured":"Sun, F., Ren, P., Yin, B., Wang, F., Li, H.: Catnet: A cascaded and aggregated transformer network for rgb-d salient object detection. IEEE Transactions on Multimedia (2023)","DOI":"10.1109\/TMM.2023.3294003"},{"key":"1717_CR21","doi-asserted-by":"crossref","unstructured":"Qin, X., Zhang, Z., Huang, C., Gao, C., Dehghan, M., Jagersand, M.: Basnet: Boundary-aware salient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 7479\u20137489 (2019)","DOI":"10.1109\/CVPR.2019.00766"},{"key":"1717_CR22","unstructured":"Howard, A.G., Zhu, M., Chen, B., Kalenichenko, D., Wang, W., Weyand, T., Andreetto, M., Adam, H.: Mobilenets: Efficient convolutional neural networks for mobile vision applications. (2017). arXiv preprint arXiv:1704.04861"},{"key":"1717_CR23","doi-asserted-by":"crossref","unstructured":"Han, K., Wang, Y., Tian, Q., Guo, J., Xu, C., Xu, C.: Ghostnet: More features from cheap operations. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1580\u20131589 (2020)","DOI":"10.1109\/CVPR42600.2020.00165"},{"key":"1717_CR24","doi-asserted-by":"crossref","unstructured":"Yang, C., Zhang, L., Lu, H., Ruan, X., Yang, M.-H.: Saliency detection via graph-based manifold ranking. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3166\u20133173 (2013)","DOI":"10.1109\/CVPR.2013.407"},{"issue":"4","key":"1717_CR25","doi-asserted-by":"publisher","first-page":"818","DOI":"10.1109\/TPAMI.2016.2562626","volume":"39","author":"H Peng","year":"2016","unstructured":"Peng, H., Li, B., Ling, H., Hu, W., Xiong, W., Maybank, S.J.: Salient object detection via structured matrix decomposition. IEEE Trans. Pattern Anal. Mach. Intell. 39(4), 818\u2013832 (2016)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"issue":"10","key":"1717_CR26","doi-asserted-by":"publisher","first-page":"3176","DOI":"10.1109\/TIP.2015.2440174","volume":"24","author":"H Li","year":"2015","unstructured":"Li, H., Lu, H., Lin, Z., Shen, X., Price, B.: Inner and inter label propagation: salient object detection in the wild. IEEE Trans. Image Process. 24(10), 3176\u20133186 (2015)","journal-title":"IEEE Trans. Image Process."},{"key":"1717_CR27","doi-asserted-by":"crossref","unstructured":"Pang, Y., Zhao, X., Zhang, L., Lu, H.: Multi-scale interactive network for salient object detection. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9413\u20139422 (2020)","DOI":"10.1109\/CVPR42600.2020.00943"},{"key":"1717_CR28","unstructured":"Yun, Y.K., Lin, W.: Selfreformer: Self-refined network with transformer for salient object detection. (2022). arXiv preprint arXiv:2205.11283"},{"key":"1717_CR29","doi-asserted-by":"crossref","unstructured":"Deng, X., Zhang, P., Liu, W., Lu, H.: Recurrent multi-scale transformer for high-resolution salient object detection. In: Proceedings of the 31st ACM International Conference on Multimedia, pp. 7413\u20137423 (2023)","DOI":"10.1145\/3581783.3611983"},{"key":"1717_CR30","doi-asserted-by":"crossref","unstructured":"Qiu, Y., Liu, Y., Zhang, L., Lu, H., Xu, J.: Boosting salient object detection with transformer-based asymmetric bilateral u-net. IEEE Transactions on Circuits and Systems for Video Technology (2023)","DOI":"10.1109\/TCSVT.2023.3307693"},{"key":"1717_CR31","doi-asserted-by":"crossref","unstructured":"Lu, X., Yuan, Y., Liu, X., Wang, L., Zhou, X., Yang, Y.: Low-light salient object detection by learning to highlight the foreground objects. IEEE Transactions on Circuits and Systems for Video Technology (2024)","DOI":"10.1109\/TCSVT.2024.3377108"},{"issue":"5","key":"1717_CR32","doi-asserted-by":"publisher","first-page":"2274","DOI":"10.1109\/TIP.2017.2682981","volume":"26","author":"L Qu","year":"2017","unstructured":"Qu, L., He, S., Zhang, J., Tian, J., Tang, Y., Yang, Q.: Rgbd salient object detection via deep fusion. IEEE Trans. Image Process. 26(5), 2274\u20132285 (2017)","journal-title":"IEEE Trans. Image Process."},{"key":"1717_CR33","doi-asserted-by":"crossref","unstructured":"Zhang, W., Ji, G.-P., Wang, Z., Fu, K., Zhao, Q.: Depth quality-inspired feature manipulation for efficient rgb-d salient object detection. In: Proceedings of the 29th ACM International Conference on Multimedia, pp. 731\u2013740 (2021)","DOI":"10.1145\/3474085.3475240"},{"key":"1717_CR34","doi-asserted-by":"publisher","first-page":"7012","DOI":"10.1109\/TIP.2020.3028289","volume":"30","author":"Z Chen","year":"2020","unstructured":"Chen, Z., Cong, R., Xu, Q., Huang, Q.: Dpanet: Depth potentiality-aware gated attention network for rgb-d salient object detection. IEEE Trans. Image Process. 30, 7012\u20137024 (2020)","journal-title":"IEEE Trans. Image Process."},{"key":"1717_CR35","doi-asserted-by":"publisher","first-page":"152","DOI":"10.1016\/j.neucom.2022.12.004","volume":"522","author":"T Chen","year":"2023","unstructured":"Chen, T., Xiao, J., Hu, X., Zhang, G., Wang, S.: Adaptive fusion network for rgb-d salient object detection. Neurocomputing 522, 152\u2013164 (2023)","journal-title":"Neurocomputing"},{"key":"1717_CR36","unstructured":"Yin, B., Zhang, X., Li, Z., Liu, L., Cheng, M.-M., Hou, Q.: Dformer: Rethinking rgbd representation learning for semantic segmentation. (2023). arXiv preprint arXiv:2309.09668"},{"key":"1717_CR37","doi-asserted-by":"publisher","first-page":"6800","DOI":"10.1109\/TIP.2022.3216198","volume":"31","author":"R Cong","year":"2022","unstructured":"Cong, R., Lin, Q., Zhang, C., Li, C., Cao, X., Huang, Q., Zhao, Y.: Cir-net: Cross-modality interaction and refinement for rgb-d salient object detection. IEEE Trans. Image Process. 31, 6800\u20136815 (2022)","journal-title":"IEEE Trans. Image Process."},{"key":"1717_CR38","doi-asserted-by":"crossref","unstructured":"Ji, W., Li, J., Zhang, M., Piao, Y., Lu, H.: Accurate rgb-d salient object detection via collaborative learning. In: Computer Vision\u2013ECCV 2020: 16th European Conference, Glasgow, UK, August 23\u201328, 2020, Proceedings, Part XVIII 16, pp. 52\u201369 (2020). Springer","DOI":"10.1007\/978-3-030-58523-5_4"},{"key":"1717_CR39","doi-asserted-by":"crossref","unstructured":"Ren, S., Zhao, N., Wen, Q., Han, G., He, S.: Unifying global-local representations in salient object detection with transformers. IEEE Transactions on Emerging Topics in Computational Intelligence (2024)","DOI":"10.1109\/TETCI.2024.3380442"},{"key":"1717_CR40","doi-asserted-by":"crossref","unstructured":"Liu, N., Zhang, N., Wan, K., Shao, L., Han, J.: Visual saliency transformer. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 4722\u20134732 (2021)","DOI":"10.1109\/ICCV48922.2021.00468"},{"key":"1717_CR41","unstructured":"Wang, X., Wang, X., Jiang, B., Tang, J., Luo, B.: Mutualformer: Multi-modality representation learning via cross-diffusion attention. (2021). arXiv preprint arXiv:2112.01177"},{"issue":"7","key":"1717_CR42","doi-asserted-by":"publisher","first-page":"4486","DOI":"10.1109\/TCSVT.2021.3127149","volume":"32","author":"Z Liu","year":"2021","unstructured":"Liu, Z., Tan, Y., He, Q., Xiao, Y.: Swinnet: Swin transformer drives edge-aware rgb-d and rgb-t salient object detection. IEEE Trans. Circ. Syst. Video Tech. 32(7), 4486\u20134497 (2021)","journal-title":"IEEE Trans. Circ. Syst. Video Tech."},{"key":"1717_CR43","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2023.126779","volume":"559","author":"C Zeng","year":"2023","unstructured":"Zeng, C., Kwong, S., Ip, H.: Dual swin-transformer based mutual interactive network for rgb-d salient object detection. Neurocomputing 559, 126779 (2023)","journal-title":"Neurocomputing"},{"key":"1717_CR44","doi-asserted-by":"crossref","unstructured":"Liu, Z., Wang, Y., Tu, Z., Xiao, Y., Tang, B.: Tritransnet: Rgb-d salient object detection with a triplet transformer embedding network. In: Proceedings of the 29th ACM International Conference on Multimedia, pp. 4481\u20134490 (2021)","DOI":"10.1145\/3474085.3475601"},{"key":"1717_CR45","doi-asserted-by":"crossref","unstructured":"Lin, T.-Y., Doll\u00e1r, P., Girshick, R., He, K., Hariharan, B., Belongie, S.: Feature pyramid networks for object detection. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 2117\u20132125 (2017)","DOI":"10.1109\/CVPR.2017.106"},{"key":"1717_CR46","doi-asserted-by":"crossref","unstructured":"Woo, S., Park, J., Lee, J.-Y., Kweon, I.S.: Cbam: Convolutional block attention module. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 3\u201319 (2018)","DOI":"10.1007\/978-3-030-01234-2_1"},{"key":"1717_CR47","doi-asserted-by":"crossref","unstructured":"Liu, J.-J., Hou, Q., Cheng, M.-M., Wang, C., Feng, J.: Improving convolutional networks with self-calibrated convolutions. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10096\u201310105 (2020)","DOI":"10.1109\/CVPR42600.2020.01011"},{"key":"1717_CR48","doi-asserted-by":"crossref","unstructured":"Ju, R., Ge, L., Geng, W., Ren, T., Wu, G.: Depth saliency based on anisotropic center-surround difference. In: 2014 IEEE International Conference on Image Processing (ICIP), pp. 1115\u20131119 (2014). IEEE","DOI":"10.1109\/ICIP.2014.7025222"},{"key":"1717_CR49","doi-asserted-by":"publisher","unstructured":"Li, N., Ye, J., Ji, Y., Ling, H., Yu, J.: Saliency detection on light field. In: 2014 IEEE Conference on Computer Vision and Pattern Recognition, pp. 2806\u20132813 (2014). https:\/\/doi.org\/10.1109\/CVPR.2014.359","DOI":"10.1109\/CVPR.2014.359"},{"key":"1717_CR50","doi-asserted-by":"crossref","unstructured":"Ren, J., Gong, X., Yu, L., Zhou, W., Ying\u00a0Yang, M.: Exploiting global priors for rgb-d saliency detection. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition Workshops, pp. 25\u201332 (2015)","DOI":"10.1109\/CVPRW.2015.7301391"},{"key":"1717_CR51","doi-asserted-by":"crossref","unstructured":"Zhu, C., Li, G.: A three-pathway psychobiological framework of salient object detection using stereoscopic technology. In: Proceedings of the IEEE International Conference on Computer Vision Workshops, pp. 3008\u20133014 (2017)","DOI":"10.1109\/ICCVW.2017.355"},{"key":"1717_CR52","doi-asserted-by":"publisher","unstructured":"Niu, Y., Geng, Y., Li, X., Liu, F.: Leveraging stereopsis for saliency analysis. In: 2012 IEEE Conference on Computer Vision and Pattern Recognition, pp. 454\u2013461 (2012). https:\/\/doi.org\/10.1109\/CVPR.2012.6247708","DOI":"10.1109\/CVPR.2012.6247708"},{"key":"1717_CR53","doi-asserted-by":"crossref","unstructured":"Chen, S., Fu, Y.: Progressively guided alternate refinement network for rgb-d salient object detection. In: European Conference on Computer Vision, pp. 520\u2013538 (2020). Springer","DOI":"10.1007\/978-3-030-58598-3_31"},{"issue":"5","key":"1717_CR54","doi-asserted-by":"publisher","first-page":"2075","DOI":"10.1109\/TNNLS.2020.2996406","volume":"32","author":"D-P Fan","year":"2020","unstructured":"Fan, D.-P., Lin, Z., Zhang, Z., Zhu, M., Cheng, M.-M.: Rethinking rgb-d salient object detection: Models, data sets, and large-scale benchmarks. IEEE Trans. Neural Netw. Learn. Syst. 32(5), 2075\u20132089 (2020)","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"key":"1717_CR55","doi-asserted-by":"publisher","unstructured":"Achanta, R., Hemami, S., Estrada, F., Susstrunk, S.: Frequency-tuned salient region detection. In: 2009 IEEE Conference on Computer Vision and Pattern Recognition, pp. 1597\u20131604 (2009). https:\/\/doi.org\/10.1109\/CVPR.2009.5206596","DOI":"10.1109\/CVPR.2009.5206596"},{"key":"1717_CR56","doi-asserted-by":"crossref","unstructured":"Fan, D.-P., Cheng, M.-M., Liu, Y., Li, T., Borji, A.: Structure-measure: A New Way to Evaluate Foreground Maps (2017)","DOI":"10.1109\/ICCV.2017.487"},{"key":"1717_CR57","doi-asserted-by":"crossref","unstructured":"Fan, D.-P., Gong, C., Cao, Y., Ren, B., Cheng, M.-M., Borji, A.: Enhanced-alignment Measure for Binary Foreground Map Evaluation (2018)","DOI":"10.24963\/ijcai.2018\/97"},{"key":"1717_CR58","doi-asserted-by":"publisher","unstructured":"Perazzi, F., Kr\u00e4henb\u00fchl, P., Pritch, Y., Hornung, A.: Saliency filters: Contrast based filtering for salient region detection. In: 2012 IEEE Conference on Computer Vision and Pattern Recognition, pp. 733\u2013740 (2012). https:\/\/doi.org\/10.1109\/CVPR.2012.6247743","DOI":"10.1109\/CVPR.2012.6247743"},{"issue":"12","key":"1717_CR59","doi-asserted-by":"publisher","first-page":"5706","DOI":"10.1109\/TIP.2015.2487833","volume":"24","author":"A Borji","year":"2015","unstructured":"Borji, A., Cheng, M.-M., Jiang, H., Li, J.: Salient object detection: A benchmark. IEEE Trans. Image Process. 24(12), 5706\u20135722 (2015)","journal-title":"IEEE Trans. Image Process."},{"key":"1717_CR60","doi-asserted-by":"crossref","unstructured":"Sun, P., Zhang, W., Wang, H., Li, S., Li, X.: Deep rgb-d saliency detection with depth-sensitive attention and automatic multi-modal fusion. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1407\u20131417 (2021)","DOI":"10.1109\/CVPR46437.2021.00146"},{"key":"1717_CR61","doi-asserted-by":"crossref","unstructured":"Chen, Q., Liu, Z., Zhang, Y., Fu, K., Zhao, Q., Du, H.: Rgb-d salient object detection via 3d convolutional neural networks. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 35, pp. 1063\u20131071 (2021)","DOI":"10.1609\/aaai.v35i2.16191"},{"key":"1717_CR62","doi-asserted-by":"crossref","unstructured":"Wu, Z., Gobichettipalayam, S., Tamadazte, B., Allibert, G., Paudel, D.P., Demonceaux, C.: Robust rgb-d fusion for saliency detection. In: 2022 International Conference on 3D Vision (3DV), pp. 403\u2013413 (2022). IEEE","DOI":"10.1109\/3DV57658.2022.00052"},{"key":"1717_CR63","doi-asserted-by":"crossref","unstructured":"Wang, Y., Zhang, Y.: Three-stage bidirectional interaction network for efficient rgb-d salient object detection. In: Proceedings of the Asian Conference on Computer Vision, pp. 3672\u20133689 (2022)","DOI":"10.1007\/978-3-031-26348-4_13"},{"issue":"11","key":"1717_CR64","doi-asserted-by":"publisher","first-page":"7632","DOI":"10.1109\/TCSVT.2022.3180274","volume":"32","author":"X Jin","year":"2022","unstructured":"Jin, X., Yi, K., Xu, J.: Moadnet: Mobile asymmetric dual-stream networks for real-time and lightweight rgb-d salient object detection. IEEE Trans. Circuits Syst. Video Technol. 32(11), 7632\u20137645 (2022). https:\/\/doi.org\/10.1109\/TCSVT.2022.3180274","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"1717_CR65","doi-asserted-by":"publisher","first-page":"1285","DOI":"10.1109\/TIP.2022.3140606","volume":"31","author":"F Wang","year":"2022","unstructured":"Wang, F., Pan, J., Xu, S., Tang, J.: Learning discriminative cross-modality features for rgb-d saliency detection. IEEE Trans. Image Proces. 31, 1285\u20131297 (2022)","journal-title":"IEEE Trans. Image Proces."},{"issue":"10","key":"1717_CR66","doi-asserted-by":"publisher","first-page":"7547","DOI":"10.1007\/s00521-021-06845-3","volume":"34","author":"T Chen","year":"2022","unstructured":"Chen, T., Hu, X., Xiao, J., Zhang, G., Wang, S.: Cfidnet: Cascaded feature interaction decoder for rgb-d salient object detection. Neural Comput. Appl. 34(10), 7547\u20137563 (2022)","journal-title":"Neural Comput. Appl."},{"key":"1717_CR67","doi-asserted-by":"crossref","unstructured":"Zhang, M., Yao, S., Hu, B., Piao, Y., Ji, W.: C$$^2$$ DFNet: Criss-cross dynamic filter network for rgb-d salient object detection. IEEE Transactions on Multimedia (2022)","DOI":"10.1109\/TMM.2022.3187856"},{"key":"1717_CR68","doi-asserted-by":"publisher","DOI":"10.1016\/j.patcog.2022.109194","volume":"136","author":"H Bi","year":"2023","unstructured":"Bi, H., Wu, R., Liu, Z., Zhu, H., Zhang, C., Xiang, T.-Z.: Cross-modal hierarchical interaction network for rgb-d salient object detection. Pattern Recogn. 136, 109194 (2023)","journal-title":"Pattern Recogn."},{"key":"1717_CR69","doi-asserted-by":"publisher","first-page":"2160","DOI":"10.1109\/TIP.2023.3263111","volume":"32","author":"Z Wu","year":"2023","unstructured":"Wu, Z., Allibert, G., Meriaudeau, F., Ma, C., Demonceaux, C.: Hidanet: Rgb-d salient object detection via hierarchical depth awareness. IEEE Trans. Image Process. 32, 2160\u20132173 (2023)","journal-title":"IEEE Trans. Image Process."},{"key":"1717_CR70","doi-asserted-by":"crossref","unstructured":"Cong, R., Liu, H., Zhang, C., Zhang, W., Zheng, F., Song, R., Kwong, S.: Point-aware interaction and cnn-induced refinement network for rgb-d salient object detection. In: Proceedings of the 31st ACM International Conference on Multimedia, pp. 406\u2013416 (2023)","DOI":"10.1145\/3581783.3611982"},{"key":"1717_CR71","doi-asserted-by":"crossref","unstructured":"WU, Z., Paudel, D.P., Fan, D.-P., Wang, J., Wang, S., Demonceaux, C., Timofte, R., Van\u00a0Gool, L.: Source-free depth for object pop-out. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV), pp. 1032\u20131042 (2023)","DOI":"10.1109\/ICCV51070.2023.00101"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-025-01717-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-025-01717-5\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-025-01717-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,4,21]],"date-time":"2025-04-21T19:36:42Z","timestamp":1745264202000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-025-01717-5"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,2,24]]},"references-count":71,"journal-issue":{"issue":"2","published-print":{"date-parts":[[2025,4]]}},"alternative-id":["1717"],"URL":"https:\/\/doi.org\/10.1007\/s00530-025-01717-5","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,2,24]]},"assertion":[{"value":"30 June 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 February 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 February 2025","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no potential or explicit financial interests in this paper and have no objections to the order of authorship.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}},{"value":"Not applicable.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical approval and consent to participate"}},{"value":"We secured written informed consent from every participant for the purpose of publication.","order":4,"name":"Ethics","group":{"name":"EthicsHeading","label":"Consent for publication"}}],"article-number":"128"}}