{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,15]],"date-time":"2026-05-15T15:45:55Z","timestamp":1778859955092,"version":"3.51.4"},"reference-count":42,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2025,2,6]],"date-time":"2025-02-06T00:00:00Z","timestamp":1738800000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,2,6]],"date-time":"2025-02-06T00:00:00Z","timestamp":1738800000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62176119"],"award-info":[{"award-number":["62176119"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100001809","name":"Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62176119"],"award-info":[{"award-number":["62176119"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100001809","name":"Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62176119"],"award-info":[{"award-number":["62176119"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100001809","name":"Natural Science Foundation of China","doi-asserted-by":"crossref","award":["62176119"],"award-info":[{"award-number":["62176119"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach Learn"],"published-print":{"date-parts":[[2025,3]]},"DOI":"10.1007\/s10994-024-06688-8","type":"journal-article","created":{"date-parts":[[2025,2,6]],"date-time":"2025-02-06T12:59:48Z","timestamp":1738846788000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["SA-CAM: Semantic-aware visual explanations for deep convolutional networks"],"prefix":"10.1007","volume":"114","author":[{"given":"Anni","family":"Yu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Qing-Long","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Lu","family":"Rao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yu-Bin","family":"Yang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,2,6]]},"reference":[{"key":"6688_CR1","unstructured":"Adebayo, J., Gilmer, J., Muelly, M., Goodfellow, I. J., Hardt, M., & Kim, B. (2018). Sanity checks for saliency maps. In: Advances in Neural Information Processing Systems (NeurIPS 2018), (pp. 9525\u20139536)."},{"key":"6688_CR2","doi-asserted-by":"publisher","unstructured":"Bansal, N., Agarwal, C., & Nguyen, A. (2020) SAM: the sensitivity of attribution methods to hyperparameters. In: 2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), (pp. 8670\u20138680). IEEE. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00870.","DOI":"10.1109\/CVPR42600.2020.00870"},{"key":"6688_CR3","doi-asserted-by":"crossref","unstructured":"Belharbi, S., Ben\u00a0Ayed, I., McCaffrey, L., & Granger, E. (2023). Tcam: Temporal class activation maps for object localization in weakly-labeled unconstrained videos. In: Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision, (pp. 137\u2013146).","DOI":"10.1109\/WACV56688.2023.00022"},{"key":"6688_CR4","doi-asserted-by":"publisher","unstructured":"Chattopadhyay, A., Sarkar, A., Howlader, P., & Balasubramanian, V. N. (2018). Grad-cam++: Generalized gradient-based visual explanations for deep convolutional networks. In: 2018 IEEE Winter Conference on Applications of Computer Vision (WACV 2018), (pp. 839\u2013847). https:\/\/doi.org\/10.1109\/WACV.2018.00097.","DOI":"10.1109\/WACV.2018.00097"},{"key":"6688_CR5","doi-asserted-by":"crossref","unstructured":"Chefer, H., Gur, S., & Wolf, L. (2021). Generic attention-model explainability for interpreting bi-modal and encoder-decoder transformers. In: 2021 IEEE\/CVF International Conference on Computer Vision (ICCV 2021), (pp. 397\u2013406).","DOI":"10.1109\/ICCV48922.2021.00045"},{"key":"6688_CR6","doi-asserted-by":"crossref","unstructured":"Chefer, H., Gur, S., & Wolf, L. (2021). Transformer interpretability beyond attention visualization. In: 2021 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR 2021), (pp. 782\u2013791).","DOI":"10.1109\/CVPR46437.2021.00084"},{"key":"6688_CR7","doi-asserted-by":"crossref","unstructured":"Chen, Z., & Sun, Q. (2023). Extracting class activation maps from non-discriminative features as well. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, (pp. 3135\u20133144).","DOI":"10.1109\/CVPR52729.2023.00306"},{"key":"6688_CR8","unstructured":"Coates, A., Ng, A., & Lee, H. (2011). An analysis of single-layer networks in unsupervised feature learning. In: Proceedings of the Fourteenth International Conference on Artificial Intelligence and Statistics, (pp. 215\u2013223). JMLR Workshop and Conference Proceedings."},{"key":"6688_CR9","unstructured":"Dabkowski, P., & Gal, Y. (2017). Real time image saliency for black box classifiers. In: Annual Conference on Neural Information Processing Systems (NeurIPS 2017), (pp. 6967\u20136976)."},{"key":"6688_CR10","doi-asserted-by":"crossref","unstructured":"Gu, J., & Dong, C. (2021). Interpreting super-resolution networks with local attribution maps. In: IEEE Conference on Computer Vision and Pattern Recognition, CVPR 2021, Virtual, June 19-25, (pp. 9199\u20139208). Computer Vision Foundation \/ IEEE.","DOI":"10.1109\/CVPR46437.2021.00908"},{"issue":"3","key":"6688_CR11","doi-asserted-by":"publisher","first-page":"328","DOI":"10.1007\/s11263-014-0713-9","volume":"110","author":"M Guillaumin","year":"2014","unstructured":"Guillaumin, M., K\u00fcttel, D., & Ferrari, V. (2014). Imagenet auto-annotation with segmentation propagation. International Journal of Computer Vision, 110(3), 328\u2013348. https:\/\/doi.org\/10.1007\/s11263-014-0713-9","journal-title":"International Journal of Computer Vision"},{"key":"6688_CR12","doi-asserted-by":"crossref","unstructured":"Guo, Z., Yan, H., Li, H., & Lin, X. (2023). Class attention transfer based knowledge distillation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, (pp. 11868\u201311877).","DOI":"10.1109\/CVPR52729.2023.01142"},{"key":"6688_CR13","doi-asserted-by":"publisher","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. (2016). Deep residual learning for image recognition. In: 2016 IEEE Conference on Computer Vision and Pattern Recognition (CVPR), (pp. 770\u2013778). IEEE Computer Society. https:\/\/doi.org\/10.1109\/CVPR.2016.90.","DOI":"10.1109\/CVPR.2016.90"},{"issue":"9","key":"6688_CR14","doi-asserted-by":"publisher","first-page":"6124","DOI":"10.1007\/s00330-023-09590-4","volume":"33","author":"Y Jun","year":"2023","unstructured":"Jun, Y., Park, Y. W., Shin, H., Shin, Y., Lee, J. R., Han, K., Ahn, S. S., Lim, S. M., Hwang, D., & Lee, S.-K. (2023). Intelligent noninvasive meningioma grading with a fully automatic segmentation using interpretable multiparametric deep learning. European Radiology, 33(9), 6124\u20136133.","journal-title":"European Radiology"},{"key":"6688_CR15","doi-asserted-by":"publisher","unstructured":"Kapishnikov, A., Bolukbasi, T., Vi\u00e9gas, F. B., & Terry, M. (2019). XRAI: better attributions through regions. In: 2019 IEEE\/CVF International Conference on Computer Vision (ICCV 2019), (pp. 4947\u20134956). https:\/\/doi.org\/10.1109\/ICCV.2019.00505.","DOI":"10.1109\/ICCV.2019.00505"},{"key":"6688_CR16","unstructured":"Lei, Y., Li, Z., Li, Y., Zhang, J., & Shan, H. (2023). Lico: Explainable models with language-image consistency. In: Thirty-seventh Conference on Neural Information Processing Systems."},{"key":"6688_CR17","doi-asserted-by":"publisher","unstructured":"Lin, T., Maire, M., Belongie, S. J., Hays, J., Perona, P., Ramanan, D., Doll\u00e1r, P., & Zitnick, C. L. (2014). Microsoft COCO: common objects in context. In: 3th European Conference on Computer Vision (ECCV). Lecture Notes in Computer Science, (vol. 8693, pp. 740\u2013755). Springer. https:\/\/doi.org\/10.1007\/978-3-319-10602-1_48.","DOI":"10.1007\/978-3-319-10602-1_48"},{"key":"6688_CR18","doi-asserted-by":"crossref","unstructured":"Luo, Z., Liu, Y., Schiele, B., & Sun, Q. (2023). Class-incremental exemplar compression for class-incremental learning. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, (pp. 11371\u201311380).","DOI":"10.1109\/CVPR52729.2023.01094"},{"issue":"5","key":"6688_CR19","doi-asserted-by":"publisher","first-page":"3503","DOI":"10.1007\/s10462-021-10088-y","volume":"55","author":"D Minh","year":"2022","unstructured":"Minh, D., Wang, H. X., Li, Y. F., & Nguyen, T. N. (2022). Explainable artificial intelligence: a comprehensive review. Artificial Intelligence Review, 55(5), 3503\u20133568.","journal-title":"Artificial Intelligence Review"},{"key":"6688_CR20","first-page":"2825","volume":"12","author":"F Pedregosa","year":"2011","unstructured":"Pedregosa, F., Varoquaux, G., Gramfort, A., Michel, V., Thirion, B., Grisel, O., Blondel, M., Prettenhofer, P., Weiss, R., Dubourg, V., Vanderplas, J., Passos, A., Cournapeau, D., Brucher, M., Perrot, M., & Duchesnay, E. (2011). Scikit-learn: Machine learning in Python. Journal of Machine Learning Research, 12, 2825\u20132830.","journal-title":"Journal of Machine Learning Research"},{"key":"6688_CR21","unstructured":"Petsiuk, V., Das, A., & Saenko, K. (2018). RISE: randomized input sampling for explanation of black-box models. In: British Machine Vision Conference (BMCV 2018), (p. 151)."},{"key":"6688_CR22","unstructured":"Radford, A., Kim, J. W., Hallacy, C., Ramesh, A., Goh, G., Agarwal, S., Sastry, G., Askell, A., Mishkin, P., & Clark, J., et al. (2021). Learning transferable visual models from natural language supervision. In: International Conference on Machine Learning, (pp. 8748\u20138763). PMLR."},{"key":"6688_CR23","doi-asserted-by":"publisher","unstructured":"Rebuffi, S., Fong, R., Ji, X., & Vedaldi, A. (2020). There and back again: Revisiting backpropagation saliency methods. In: 2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR 2020), Seattle, WA, USA, June 13-19, (pp. 8836\u20138845). IEEE. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00886.","DOI":"10.1109\/CVPR42600.2020.00886"},{"issue":"3","key":"6688_CR24","doi-asserted-by":"publisher","first-page":"211","DOI":"10.1007\/s11263-015-0816-y","volume":"115","author":"O Russakovsky","year":"2015","unstructured":"Russakovsky, O., Deng, J., Su, H., Krause, J., Satheesh, S., Ma, S., Huang, Z., Karpathy, A., Khosla, A., Bernstein, M. S., Berg, A. C., & Li, F. (2015). Imagenet large scale visual recognition challenge. International Journal of Computer Vision, 115(3), 211\u2013252. https:\/\/doi.org\/10.1007\/s11263-015-0816-y","journal-title":"International Journal of Computer Vision"},{"issue":"11","key":"6688_CR25","doi-asserted-by":"publisher","first-page":"2660","DOI":"10.1109\/TNNLS.2016.2599820","volume":"28","author":"W Samek","year":"2017","unstructured":"Samek, W., Binder, A., Montavon, G., Lapuschkin, S., & M\u00fcller, K. (2017). Evaluating the visualization of what a deep neural network has learned. IEEE Trans. Neural Networks Learn. Syst., 28(11), 2660\u20132673. https:\/\/doi.org\/10.1109\/TNNLS.2016.2599820","journal-title":"IEEE Trans. Neural Networks Learn. Syst."},{"key":"6688_CR26","doi-asserted-by":"publisher","unstructured":"Selvaraju, R. R., Cogswell, M., Das, A., Vedantam, R., Parikh, D., & Batra, D. (2017). Grad-cam: Visual explanations from deep networks via gradient-based localization. In: IEEE International Conference on Computer Vision (ICCV 2017), (pp. 618\u2013626). https:\/\/doi.org\/10.1109\/ICCV.2017.74.","DOI":"10.1109\/ICCV.2017.74"},{"key":"6688_CR27","doi-asserted-by":"publisher","first-page":"8572","DOI":"10.1109\/ACCESS.2019.2963055","volume":"8","author":"D Seo","year":"2020","unstructured":"Seo, D., Oh, K., & Oh, I. (2020). Regional multi-scale approach for visually pleasing explanations of deep neural networks. IEEE Access, 8, 8572\u20138582. https:\/\/doi.org\/10.1109\/ACCESS.2019.2963055","journal-title":"IEEE Access"},{"key":"6688_CR28","unstructured":"Simonyan, K., & Zisserman, A. (2015). Very deep convolutional networks for large-scale image recognition. In: 3rd International Conference on Learning Representations (ICLR 2015)."},{"key":"6688_CR29","unstructured":"Smilkov, D., Thorat, N., Kim, B., Vi\u00e9gas, F. B., & Wattenberg, M. (2017). Smoothgrad: removing noise by adding noise. arXiv preprint arXiv:1706.03825."},{"issue":"4","key":"6688_CR30","doi-asserted-by":"publisher","first-page":"1839","DOI":"10.1007\/s10994-022-06245-1","volume":"113","author":"P Tan","year":"2024","unstructured":"Tan, P., Tan, Z.-H., Jiang, Y., & Zhou, Z.-H. (2024). Towards enabling learnware to handle heterogeneous feature spaces. Machine Learning, 113(4), 1839\u20131860.","journal-title":"Machine Learning"},{"key":"6688_CR31","doi-asserted-by":"publisher","unstructured":"Wang, H., Wang, Z., Du, M., Yang, F., Zhang, Z., Ding, S., Mardziel, P., & Hu, X. (2020). Score-cam: Score-weighted visual explanations for convolutional neural networks. In: CVPR Workshops 2020, (pp. 111\u2013119). https:\/\/doi.org\/10.1109\/CVPRW50498.2020.00020.","DOI":"10.1109\/CVPRW50498.2020.00020"},{"key":"6688_CR32","doi-asserted-by":"crossref","unstructured":"Wang, H., Wang, Z., Du, M., Yang, F., Zhang, Z., Ding, S., Mardziel, P., & Hu, X. (2020). Score-cam: Score-weighted visual explanations for convolutional neural networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops, (pp. 24\u201325).","DOI":"10.1109\/CVPRW50498.2020.00020"},{"key":"6688_CR33","doi-asserted-by":"publisher","first-page":"6701","DOI":"10.1109\/TIP.2021.3097187","volume":"30","author":"D Wang","year":"2021","unstructured":"Wang, D., Cui, X., Chen, X., Ward, R., & Wang, Z. J. (2021). Interpreting bottom-up decision-making of cnns via hierarchical inference. IEEE Transactions on Image Processing, 30, 6701\u20136714. https:\/\/doi.org\/10.1109\/TIP.2021.3097187","journal-title":"IEEE Transactions on Image Processing"},{"key":"6688_CR34","unstructured":"Wu, J.-H., Zhang, S.-Q., Jiang, Y., & Zhou, Z.-H. (2024). Complex-valued neurons can learn more but slower than real-valued neurons via gradient descent. Advances in Neural Information Processing Systems36."},{"key":"6688_CR35","doi-asserted-by":"crossref","unstructured":"Xu, L., Bennamoun, M., Boussaid, F., Laga, H., Ouyang, W., & Xu, D. (2024). Mctformer+: Multi-class token transformer for weakly supervised semantic segmentation. IEEE Transactions on Pattern Analysis and Machine Intelligence.","DOI":"10.1109\/TPAMI.2024.3404422"},{"key":"6688_CR36","doi-asserted-by":"crossref","unstructured":"Xu, L., Ouyang, W., Bennamoun, M., Boussaid, F., & Xu, D. (2023). Learning multi-modal class-specific tokens for weakly supervised dense object localization. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, (pp. 19596\u201319605).","DOI":"10.1109\/CVPR52729.2023.01877"},{"key":"6688_CR37","doi-asserted-by":"publisher","unstructured":"Xu, S., Venugopalan, S., & Sundararajan, M. (2020). Attribution in scale and space. In: 2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), (pp. 9677\u20139686). IEEE. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00970.","DOI":"10.1109\/CVPR42600.2020.00970"},{"key":"6688_CR38","unstructured":"Zhang, H., Ciss\u00e9, M., Dauphin, Y. N., & Lopez-Paz, D. (2018). mixup: Beyond empirical risk minimization. In: 6th International Conference on Learning Representations, ICLR 2018, Vancouver, BC, Canada, April 30 - May 3, Conference Track Proceedings. OpenReview.net. https:\/\/openreview.net\/forum?id=r1Ddp1-Rb."},{"key":"6688_CR39","doi-asserted-by":"crossref","unstructured":"Zhang, Q., Rao, L., & Yang, Y. (2021). A novel visual interpretability for deep neural networks by optimizing activation maps with perturbation. In: Thirty-Fifth AAAI Conference on Artificial Intelligence (AAAI), (pp. 3377\u20133384). AAAI Press.","DOI":"10.1609\/aaai.v35i4.16450"},{"issue":"10","key":"6688_CR40","doi-asserted-by":"publisher","first-page":"1084","DOI":"10.1007\/s11263-017-1059-x","volume":"126","author":"J Zhang","year":"2018","unstructured":"Zhang, J., Bargal, S. A., Lin, Z., Brandt, J., Shen, X., & Sclaroff, S. (2018). Top-down neural attention by excitation backprop. International Journal of Computer Vision, 126(10), 1084\u20131102. https:\/\/doi.org\/10.1007\/s11263-017-1059-x","journal-title":"International Journal of Computer Vision"},{"key":"6688_CR41","doi-asserted-by":"publisher","unstructured":"Zhou, B., Khosla, A., Lapedriza, \u00c0., Oliva, A., & Torralba, A. (2016). Learning deep features for discriminative localization. In: 2016 IEEE Conference on Computer Vision and Pattern Recognition (CVPR 2016), (pp. 2921\u20132929). https:\/\/doi.org\/10.1109\/CVPR.2016.319.","DOI":"10.1109\/CVPR.2016.319"},{"key":"6688_CR42","doi-asserted-by":"crossref","unstructured":"Zhu, L., She, Q., Chen, Q., Meng, X., Geng, M., Jin, L., Zhang, Y., Ren, Q., & Lu, Y. (2023). Background-aware classification activation map for weakly supervised object localization. IEEE Transactions on Pattern Analysis and Machine Intelligence.","DOI":"10.1109\/TPAMI.2023.3309621"}],"container-title":["Machine Learning"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-024-06688-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10994-024-06688-8","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-024-06688-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,2,6]],"date-time":"2026-02-06T01:02:11Z","timestamp":1770339731000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10994-024-06688-8"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,2,6]]},"references-count":42,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,3]]}},"alternative-id":["6688"],"URL":"https:\/\/doi.org\/10.1007\/s10994-024-06688-8","relation":{},"ISSN":["0885-6125","1573-0565"],"issn-type":[{"value":"0885-6125","type":"print"},{"value":"1573-0565","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,2,6]]},"assertion":[{"value":"29 May 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"9 August 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"12 December 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 February 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"This content has been made available to all.","name":"free","label":"Free to read"}],"article-number":"53"}}