{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T17:04:00Z","timestamp":1784567040638,"version":"3.55.0"},"reference-count":95,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2026,5,10]],"date-time":"2026-05-10T00:00:00Z","timestamp":1778371200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,10]],"date-time":"2026-05-10T00:00:00Z","timestamp":1778371200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62405255"],"award-info":[{"award-number":["62405255"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100021171","name":"Basic and Applied Basic Research Foundation of Guangdong Province","doi-asserted-by":"publisher","award":["2023A1515110679"],"award-info":[{"award-number":["2023A1515110679"]}],"id":[{"id":"10.13039\/501100021171","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s11263-026-02872-6","type":"journal-article","created":{"date-parts":[[2026,5,10]],"date-time":"2026-05-10T08:57:15Z","timestamp":1778403435000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Dynamic Token Masking in Spiking Neural Network"],"prefix":"10.1007","volume":"134","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-0228-9082","authenticated-orcid":false,"given":"Yuetong","family":"Fang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0004-8940-0461","authenticated-orcid":false,"given":"Ziqing","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Deming","family":"Zhou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Hongwei","family":"Ren","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Shibo","family":"Zhou","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Renjing","family":"Xu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,10]]},"reference":[{"key":"2872_CR1","unstructured":"Achiam, J., Adler, S., Agarwal, S., Ahmad, L., Akkaya, I., Aleman, F.L., Almeida, D., Altenschmidt, J., Altman, S., Anadkat, S., & et al. Gpt-4 technical report. arXiv:2303.08774 (2023)"},{"issue":"10","key":"2872_CR2","doi-asserted-by":"publisher","first-page":"1537","DOI":"10.1109\/TCAD.2015.2474396","volume":"34","author":"F Akopyan","year":"2015","unstructured":"Akopyan, F., Sawada, J., Cassidy, A., Alvarez-Icaza, R., Arthur, J., Merolla, P., Imam, N., Nakamura, Y., Datta, P., & Nam, G.-J. (2015). Truenorth: Design and tool flow of a 65 mw 1 million neuron programmable neurosynaptic chip. IEEE transactions on computer-aided design of integrated circuits and systems, 34(10), 1537\u20131557.","journal-title":"IEEE transactions on computer-aided design of integrated circuits and systems"},{"key":"2872_CR3","doi-asserted-by":"crossref","unstructured":"Basu, A., Deng, L., Frenkel, C., & Zhang, X. Spiking neural network integrated circuits: A review of trends and future directions. In: 2022 IEEE Custom Integrated Circuits Conference (CICC), pp. 1\u20138 (2022). IEEE","DOI":"10.1109\/CICC53496.2022.9772783"},{"key":"2872_CR4","doi-asserted-by":"crossref","unstructured":"Bi, Y., Chadha, A., Abbas, A., Bourtsoulatze, E., & Andreopoulos, Y. Graph-based object classification for neuromorphic vision sensing. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 491\u2013501 (2019)","DOI":"10.1109\/ICCV.2019.00058"},{"key":"2872_CR5","doi-asserted-by":"crossref","unstructured":"Bos, H., & Muir, D. Sub-mw neuromorphic snn audio processing applications with rockpool and xylo. In: Embedded Artificial Intelligence, pp. 69\u201378 (2023)","DOI":"10.1201\/9781003394440-7"},{"key":"2872_CR6","doi-asserted-by":"crossref","unstructured":"Bu, T., Ding, J., Yu, Z., & Huang, T. Optimized potential initialization for low-latency spiking neural networks. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 36, pp. 11\u201320 (2022)","DOI":"10.1609\/aaai.v36i1.19874"},{"key":"2872_CR7","unstructured":"Bu, T., Fang, W., Ding, J., Dai, P., Yu, Z., & Huang, T. Optimal ANN-SNN Conversion for High-accuracy and Ultra-low-latency Spiking Neural Networks. In: International Conference on Learning Representations (2021)"},{"key":"2872_CR8","doi-asserted-by":"publisher","first-page":"54","DOI":"10.1007\/s11263-014-0788-3","volume":"113","author":"Y Cao","year":"2015","unstructured":"Cao, Y., Chen, Y., & Khosla, D. (2015). Spiking deep convolutional neural networks for energy-efficient object recognition. International Journal of Computer Vision, 113, 54\u201366.","journal-title":"International Journal of Computer Vision"},{"key":"2872_CR9","unstructured":"Choromanski, K., Likhosherstov, V., Dohan, D., Song, X., Gane, A., Sarlos, T., Hawkins, P., Davis, J., Mohiuddin, A., Kaiser, L., & et al. Rethinking attention with performers. arXiv:2009.14794 (2020)"},{"key":"2872_CR10","first-page":"9355","volume":"34","author":"X Chu","year":"2021","unstructured":"Chu, X., Tian, Z., Wang, Y., Zhang, B., Ren, H., Wei, X., Xia, H., & Shen, C. (2021). Twins: Revisiting the design of spatial attention in vision transformers. Advances in neural information processing systems, 34, 9355\u20139366.","journal-title":"Advances in neural information processing systems"},{"issue":"18","key":"2872_CR11","doi-asserted-by":"publisher","first-page":"921","DOI":"10.1016\/j.cub.2014.08.026","volume":"24","author":"DD Cox","year":"2014","unstructured":"Cox, D. D., & Dean, T. (2014). Neural networks and neuroscience-inspired computer vision. Current Biology, 24(18), 921\u2013929.","journal-title":"Current Biology"},{"issue":"1","key":"2872_CR12","doi-asserted-by":"publisher","first-page":"82","DOI":"10.1109\/MM.2018.112130359","volume":"38","author":"M Davies","year":"2018","unstructured":"Davies, M., Srinivasa, N., Lin, T.-H., Chinya, G., Cao, Y., Choday, S. H., Dimou, G., Joshi, P., Imam, N., Jain, S., et al. (2018). Loihi: A neuromorphic manycore processor with on-chip learning. Ieee Micro, 38(1), 82\u201399.","journal-title":"Ieee Micro"},{"issue":"5","key":"2872_CR13","doi-asserted-by":"publisher","first-page":"911","DOI":"10.1109\/JPROC.2021.3067593","volume":"109","author":"M Davies","year":"2021","unstructured":"Davies, M., Wild, A., Orchard, G., Sandamirskaya, Y., Guerra, G. A. F., Joshi, P., Plank, P., & Risbud, S. R. (2021). Advancing neuromorphic computing with loihi: A survey of results and outlook. Proceedings of the IEEE, 109(5), 911\u2013934.","journal-title":"Proceedings of the IEEE"},{"key":"2872_CR14","doi-asserted-by":"crossref","unstructured":"Deng, J., Dong, W., Socher, R., Li, L.-J., Li, K., & Fei-Fei, L. Imagenet: A large-scale hierarchical image database. In: 2009 IEEE Conference on Computer Vision and Pattern Recognition, pp. 248\u2013255 (2009)","DOI":"10.1109\/CVPR.2009.5206848"},{"key":"2872_CR15","unstructured":"Deng, S., Li, Y., Zhang, S., & Gu, S. Temporal Efficient Training of Spiking Neural Network via Gradient Re-weighting. arXiv:2202.11946 (2022)"},{"issue":"1","key":"2872_CR16","doi-asserted-by":"publisher","first-page":"193","DOI":"10.1146\/annurev.ne.18.030195.001205","volume":"18","author":"R Desimone","year":"1995","unstructured":"Desimone, R., Duncan, J., et al. (1995). Neural mechanisms of selective visual attention. Annual review of neuroscience, 18(1), 193\u2013222.","journal-title":"Annual review of neuroscience"},{"key":"2872_CR17","doi-asserted-by":"crossref","unstructured":"Ding, J., Yu, Z., Tian, Y., & Huang, T. Optimal ann-snn conversion for fast and accurate inference in deep spiking neural networks. arXiv:2105.11654 (2021)","DOI":"10.24963\/ijcai.2021\/321"},{"key":"2872_CR18","unstructured":"Fang, W., Chen, Y., Ding, J., Chen, D., Yu, Z., Zhou, H., Masquelier, T., & Tian, Y., contributors: SpikingJelly. https:\/\/github.com\/fangwei123456\/spikingjelly. Accessed: 2023-02-21 (2020)"},{"key":"2872_CR19","doi-asserted-by":"crossref","unstructured":"Fang, Y., Wang, Z., Zhang, L., Cao, J., Chen, H., & Xu, R. Spiking wavelet transformer. arXiv:2403.11138 (2024)","DOI":"10.1007\/978-3-031-73116-7_2"},{"key":"2872_CR20","doi-asserted-by":"crossref","unstructured":"Fang, W., Yu, Z., Chen, Y., Masquelier, T., Huang, T., & Tian, Y. Incorporating learnable membrane time constant to enhance learning of spiking neural networks. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 2661\u20132671 (2021)","DOI":"10.1109\/ICCV48922.2021.00266"},{"key":"2872_CR21","unstructured":"Fang, Y., Zhou, D., Wang, Z., Ren, H., Zeng, Z., Li, L., Zhou, S., & Xu, R. Spiking transformers need high frequency information. arXiv:2505.18608 (2025)"},{"key":"2872_CR22","first-page":"21056","volume":"34","author":"W Fang","year":"2021","unstructured":"Fang, W., Yu, Z., Chen, Y., Huang, T., Masquelier, T., & Tian, Y. (2021). Deep residual learning in spiking neural networks. Advances in Neural Information Processing Systems, 34, 21056\u201321069.","journal-title":"Advances in Neural Information Processing Systems"},{"issue":"2","key":"2872_CR23","doi-asserted-by":"publisher","first-page":"385","DOI":"10.1016\/S0896-6273(00)80788-6","volume":"23","author":"Z Gil","year":"1999","unstructured":"Gil, Z., Connors, B. W., & Amitai, Y. (1999). Efficacy of thalamocortical and intracortical synaptic connections: quanta, innervation, and reliability. Neuron, 23(2), 385\u2013397.","journal-title":"Neuron"},{"key":"2872_CR24","unstructured":"Guibas, J., Mardani, M., Li, Z., Tao, A., Anandkumar, A., & Catanzaro, B. Efficient token mixing for transformers via adaptive fourier neural operators. In: International Conference on Learning Representations (2021)"},{"key":"2872_CR25","doi-asserted-by":"crossref","unstructured":"Hao, Z., Bu, T., Ding, J., Huang, T., & Yu, Z. Reducing ann-snn conversion error through residual membrane potential. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 37, pp. 11\u201321 (2023)","DOI":"10.1609\/aaai.v37i1.25071"},{"key":"2872_CR26","doi-asserted-by":"crossref","unstructured":"He, K., Gkioxari, G., Doll\u00e1r, P., & Girshick, R. Mask r-cnn. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 2961\u20132969 (2017)","DOI":"10.1109\/ICCV.2017.322"},{"key":"2872_CR27","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778 (2016)","DOI":"10.1109\/CVPR.2016.90"},{"key":"2872_CR28","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. Identity mappings in deep residual networks. In: Computer Vision\u2013ECCV 2016: 14th European Conference, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part IV 14, pp. 630\u2013645 (2016). Springer","DOI":"10.1007\/978-3-319-46493-0_38"},{"issue":"6455","key":"2872_CR29","doi-asserted-by":"publisher","first-page":"569","DOI":"10.1038\/366569a0","volume":"366","author":"NA Hessler","year":"1993","unstructured":"Hessler, N. A., Shirke, A. M., & Malinow, R. (1993). The probability of transmitter release at a mammalian central synapse. Nature, 366(6455), 569\u2013572.","journal-title":"Nature"},{"key":"2872_CR30","doi-asserted-by":"crossref","unstructured":"Hu, Y., Deng, L., Wu, Y., Yao, M., & Li, G. Advancing spiking neural networks toward deep residual learning. IEEE Transactions on Neural Networks and Learning Systems (2024)","DOI":"10.1109\/TNNLS.2024.3355393"},{"key":"2872_CR31","doi-asserted-by":"crossref","unstructured":"Hu, Y., Zheng, Q., Jiang, X., & Pan, G. Fast-snn: fast spiking neural network by converting quantized ann. IEEE Transactions on Pattern Analysis and Machine Intelligence (2023)","DOI":"10.1109\/TPAMI.2023.3275769"},{"key":"2872_CR32","unstructured":"Hwang, S., Lee, S., Park, D., Lee, D.,& Kung, J. Spikedattention: Training-free and fully spike-driven transformer-to-snn conversion with winner-oriented spike shift for softmax operation. In: The Thirty-eighth Annual Conference on Neural Information Processing Systems"},{"key":"2872_CR33","unstructured":"Jiang, Y., Hu, K., Zhang, T., Gao, H., Liu, Y., Fang, Y., & Chen, F. Spatio-temporal approximation: A training-free snn conversion for transformers. In: The Twelfth International Conference on Learning Representations (2024)"},{"key":"2872_CR34","doi-asserted-by":"crossref","unstructured":"Jumper, J., Evans, R., Pritzel, A., Green, T., Figurnov, M., Ronneberger, O., Tunyasuvunakool, K., Bates, R., \u017d\u00eddek, A., Potapenko, A., & et al. Highly accurate protein structure prediction with alphafold. nature596(7873), 583\u2013589 (2021)","DOI":"10.1038\/s41586-021-03819-2"},{"key":"2872_CR35","unstructured":"Katz, B., & Katz, B. Nerve, Muscle, and Synapse, (1966)"},{"key":"2872_CR36","doi-asserted-by":"crossref","unstructured":"Kim, S., Park, S., Na, B., & Yoon, S. Spiking-yolo: spiking neural network for energy-efficient object detection. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 34, pp. 11270\u201311277 (2020)","DOI":"10.1609\/aaai.v34i07.6787"},{"key":"2872_CR37","doi-asserted-by":"publisher","first-page":"686","DOI":"10.1016\/j.neunet.2021.09.022","volume":"144","author":"Y Kim","year":"2021","unstructured":"Kim, Y., & Panda, P. (2021). Optimizing deeper spiking neural networks for dynamic vision sensing. Neural Networks, 144, 686\u2013698.","journal-title":"Neural Networks"},{"key":"2872_CR38","doi-asserted-by":"crossref","unstructured":"Kirillov, A., Girshick, R., He, K., & Doll\u00e1r, P. Panoptic feature pyramid networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 6399\u20136408 (2019)","DOI":"10.1109\/CVPR.2019.00656"},{"key":"2872_CR39","unstructured":"Krizhevsky, A., & Hinton, G. Learning multiple layers of features from tiny images (2009)"},{"key":"2872_CR40","doi-asserted-by":"publisher","unstructured":"Lecun, Y., Bottou, L., Bengio, Y., & Haffner, P. (1998). Gradient-based learning applied to document recognition. Proceedings of the IEEE,86(11), 2278\u20132324. https:\/\/doi.org\/10.1109\/5.726791","DOI":"10.1109\/5.726791"},{"key":"2872_CR41","doi-asserted-by":"publisher","unstructured":"Levy, W. B., & Baxter, R. A. (2002). Energy-Efficient Neuronal Computation via Quantal Synaptic Failures. Journal of Neuroscience,22(11), 4746\u20134755. https:\/\/doi.org\/10.1523\/JNEUROSCI.22-11-04746.2002. Chap. ARTICLE","DOI":"10.1523\/JNEUROSCI.22-11-04746.2002"},{"key":"2872_CR42","doi-asserted-by":"crossref","unstructured":"Li, Y., Deng, S., Dong, X., & Gu, S. Error-aware conversion from ann to snn via post-training parameter calibration. International Journal of Computer Vision, 1\u201324 (2024)","DOI":"10.1007\/s11263-024-02046-2"},{"key":"2872_CR43","unstructured":"Li, Y., Deng, S., Dong, X., Gong, R., & Gu, S. A free lunch from ANN: Towards efficient, accurate spiking neural networks calibration. In: International Conference on Machine Learning. Proceedings of Machine Learning Research, vol. 139, pp. 6316\u20136325 (2021)"},{"key":"2872_CR44","doi-asserted-by":"crossref","unstructured":"Li, Y., Kim, Y., Park, H., Geller, T., & Panda, P. Neuromorphic Data Augmentation for Training Spiking Neural Networks. arXiv:2203.06145 (2022)","DOI":"10.1007\/978-3-031-20071-7_37"},{"key":"2872_CR45","doi-asserted-by":"crossref","unstructured":"Li, H., Liu, H., Ji, X., Li, G., & Shi, L. CIFAR10-DVS: An Event-Stream Dataset for Object Classification. Frontiers in Neuroscience 11 (2017)","DOI":"10.3389\/fnins.2017.00309"},{"key":"2872_CR46","first-page":"23426","volume":"34","author":"Y Li","year":"2021","unstructured":"Li, Y., Guo, Y., Zhang, S., Deng, S., Hai, Y., & Gu, S. (2021). Differentiable spike: Rethinking gradient-descent for training spiking neural networks. Advances in Neural Information Processing Systems, 34, 23426\u201323439.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"2872_CR47","doi-asserted-by":"crossref","unstructured":"Lin, T.-Y., Maire, M., Belongie, S., Hays, J., Perona, P., Ramanan, D., Doll\u00e1r, P., & Zitnick, C.L. Microsoft coco: Common objects in context. In: Computer Vision\u2013ECCV 2014: 13th European Conference, Zurich, Switzerland, September 6-12, 2014, Proceedings, Part V 13, pp. 740\u2013755 (2014). Springer","DOI":"10.1007\/978-3-319-10602-1_48"},{"key":"2872_CR48","unstructured":"Liu, S., Li, F., Zhang, H., Yang, X., Qi, X., Su, H., Zhu, J., & Zhang, L. Dab-detr: Dynamic anchor boxes are better queries for detr. arXiv:2201.12329 (2022)"},{"key":"2872_CR49","doi-asserted-by":"crossref","unstructured":"Liu, Z., Lin, Y., Cao, Y., Hu, H., Wei, Y., Zhang, Z., Lin, S., & Guo, B. Swin transformer: Hierarchical vision transformer using shifted windows. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 10012\u201310022 (2021)","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"2872_CR50","doi-asserted-by":"crossref","unstructured":"Liu, Z., Mao, H., Wu, C.-Y., Feichtenhofer, C., Darrell, T., & Xie, S. A convnet for the 2020s. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11976\u201311986 (2022)","DOI":"10.1109\/CVPR52688.2022.01167"},{"key":"2872_CR51","doi-asserted-by":"crossref","unstructured":"Long, J., Shelhamer, E., & Darrell, T. Fully convolutional networks for semantic segmentation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3431\u20133440 (2015)","DOI":"10.1109\/CVPR.2015.7298965"},{"issue":"1","key":"2872_CR52","doi-asserted-by":"publisher","first-page":"373","DOI":"10.1146\/annurev-vision-082114-035431","volume":"1","author":"JH Maunsell","year":"2015","unstructured":"Maunsell, J. H. (2015). Neuronal mechanisms of visual attention. Annual review of vision science, 1(1), 373\u2013391.","journal-title":"Annual review of vision science"},{"key":"2872_CR53","doi-asserted-by":"crossref","unstructured":"Meng, Q., Xiao, M., Yan, S., Wang, Y., Lin, Z., & Luo, Z.-Q. Training High-Performance Low-Latency Spiking Neural Networks by Differentiation on Spike Representation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12444\u201312453 (2022)","DOI":"10.1109\/CVPR52688.2022.01212"},{"key":"2872_CR54","unstructured":"Michel, P., Levy, O., & Neubig, G. Are sixteen heads really better than one? Advances in neural information processing systems 32 (2019)"},{"key":"2872_CR55","unstructured":"Milakov, M., DevTech, S.H., & Engineer, N. Deep Learning With GPUs. Nvidia (2014)"},{"issue":"1","key":"2872_CR56","doi-asserted-by":"publisher","first-page":"47","DOI":"10.1146\/annurev-psych-122414-033400","volume":"68","author":"T Moore","year":"2017","unstructured":"Moore, T., & Zirnsak, M. (2017). Neural mechanisms of selective visual attention. Annual review of psychology, 68(1), 47\u201372.","journal-title":"Annual review of psychology"},{"issue":"6","key":"2872_CR57","doi-asserted-by":"publisher","first-page":"51","DOI":"10.1109\/MSP.2019.2931595","volume":"36","author":"EO Neftci","year":"2019","unstructured":"Neftci, E. O., Mostafa, H., & Zenke, F. (2019). Surrogate Gradient Learning in Spiking Neural Networks: Bringing the Power of Gradient-Based Optimization to Spiking Neural Networks. IEEE Signal Processing Magazine, 36(6), 51\u201363. https:\/\/doi.org\/10.1109\/MSP.2019.2931595","journal-title":"IEEE Signal Processing Magazine"},{"key":"2872_CR58","doi-asserted-by":"crossref","unstructured":"Orchard, G., Jayawant, A., & Cohen, G.K., Thakor, N. Converting Static Image Datasets to Spiking Neuromorphic Datasets Using Saccades. Frontiers in Neuroscience 9 (2015)","DOI":"10.3389\/fnins.2015.00437"},{"issue":"7767","key":"2872_CR59","doi-asserted-by":"publisher","first-page":"106","DOI":"10.1038\/s41586-019-1424-8","volume":"572","author":"J Pei","year":"2019","unstructured":"Pei, J., Deng, L., Song, S., Zhao, M., Zhang, Y., Wu, S., Wang, G., Zou, Z., Wu, Z., He, W., et al. (2019). Towards artificial general intelligence with hybrid tianjic chip architecture. Nature, 572(7767), 106\u2013111.","journal-title":"Nature"},{"key":"2872_CR60","first-page":"13937","volume":"34","author":"Y Rao","year":"2021","unstructured":"Rao, Y., Zhao, W., Liu, B., Lu, J., Zhou, J., & Hsieh, C.-J. (2021). Dynamicvit: Efficient vision transformers with dynamic token sparsification. Advances in neural information processing systems, 34, 13937\u201313949.","journal-title":"Advances in neural information processing systems"},{"key":"2872_CR61","unstructured":"Richter, O.J., Ning, Q., Liu, Q., & Sheik, S.U.A. Event-driven spiking convolutional neural network. Google Patents. US Patent App. 17\/601,939 (2022)"},{"issue":"7784","key":"2872_CR62","doi-asserted-by":"publisher","first-page":"607","DOI":"10.1038\/s41586-019-1677-2","volume":"575","author":"K Roy","year":"2019","unstructured":"Roy, K., Jaiswal, A., & Panda, P. (2019). Towards spike-based machine intelligence with neuromorphic computing. Nature, 575(7784), 607\u2013617.","journal-title":"Nature"},{"issue":"5","key":"2872_CR63","doi-asserted-by":"publisher","first-page":"353","DOI":"10.1038\/s41573-019-0050-3","volume":"19","author":"P Schneider","year":"2020","unstructured":"Schneider, P., Walters, W. P., Plowright, A. T., Sieroka, N., Listgarten, J., Goodnow, R. A., Jr., Fisher, J., Jansen, J. M., Duca, J. S., Rush, T. S., et al. (2020). Rethinking drug design in the artificial intelligence era. Nature reviews drug discovery, 19(5), 353\u2013364.","journal-title":"Nature reviews drug discovery"},{"issue":"1","key":"2872_CR64","doi-asserted-by":"publisher","first-page":"10","DOI":"10.1038\/s43588-021-00184-y","volume":"2","author":"CD Schuman","year":"2022","unstructured":"Schuman, C. D., Kulkarni, S. R., Parsa, M., Mitchell, J. P., Kay, B., et al. (2022). Opportunities for neuromorphic computing algorithms and applications. Nature Computational Science, 2(1), 10\u201319.","journal-title":"Nature Computational Science"},{"key":"2872_CR65","doi-asserted-by":"publisher","DOI":"10.1016\/j.physd.2019.132306","volume":"404","author":"A Sherstinsky","year":"2020","unstructured":"Sherstinsky, A. (2020). Fundamentals of recurrent neural network (rnn) and long short-term memory (lstm) network. Physica D: Nonlinear Phenomena, 404, Article 132306.","journal-title":"Physica D: Nonlinear Phenomena"},{"key":"2872_CR66","doi-asserted-by":"crossref","unstructured":"Shi, X., Hao, Z., & Yu, Z. Spikingresformer: Bridging resnet and vision transformer in spiking neural networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 5610\u20135619 (2024)","DOI":"10.1109\/CVPR52733.2024.00536"},{"key":"2872_CR67","doi-asserted-by":"crossref","unstructured":"Sironi, A., Brambilla, M., Bourdis, N., Lagorce, X., & Benosman, R. HATS: Histograms of Averaged Time Surfaces for Robust Event-Based Object Classification. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 1731\u20131740 (2018)","DOI":"10.1109\/CVPR.2018.00186"},{"issue":"4","key":"2872_CR68","doi-asserted-by":"publisher","DOI":"10.1088\/2634-4386\/ac8828","volume":"2","author":"KM Stewart","year":"2022","unstructured":"Stewart, K. M., & Neftci, E. O. (2022). Meta-learning spiking neural networks with surrogate gradient descent. Neuromorphic Computing and Engineering, 2(4), Article 044002.","journal-title":"Neuromorphic Computing and Engineering"},{"key":"2872_CR69","doi-asserted-by":"crossref","unstructured":"Su, Q., Chou, Y., Hu, Y., Li, J., Mei, S., Zhang, Z., & Li, G. Deep directly-trained spiking neural networks for object detection. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 6555\u20136565 (2023)","DOI":"10.1109\/ICCV51070.2023.00603"},{"key":"2872_CR70","first-page":"24261","volume":"34","author":"IO Tolstikhin","year":"2021","unstructured":"Tolstikhin, I. O., Houlsby, N., Kolesnikov, A., Beyer, L., Zhai, X., Unterthiner, T., Yung, J., Steiner, A., Keysers, D., Uszkoreit, J., et al. (2021). Mlp-mixer: An all-mlp architecture for vision. Advances in neural information processing systems, 34, 24261\u201324272.","journal-title":"Advances in neural information processing systems"},{"key":"2872_CR71","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, & Polosukhin, I. Attention is all you need. Advances in neural information processing systems 30 (2017)"},{"key":"2872_CR72","doi-asserted-by":"crossref","unstructured":"Viale, A., Marchisio, A., Martina, M., Masera, G., & Shafique, M. Carsnn: An efficient spiking neural network for event-based autonomous cars on the loihi neuromorphic research processor. In: 2021 International Joint Conference on Neural Networks (IJCNN), pp. 1\u201310 (2021)","DOI":"10.1109\/IJCNN52387.2021.9533738"},{"key":"2872_CR73","doi-asserted-by":"crossref","unstructured":"Voita, E., Talbot, D., Moiseev, F., Sennrich, R.,& Titov, I. Analyzing multi-head self-attention: Specialized heads do the heavy lifting, the rest can be pruned. arXiv:1905.09418 (2019)","DOI":"10.18653\/v1\/P19-1580"},{"key":"2872_CR74","doi-asserted-by":"crossref","unstructured":"Wang, Z., Fang, Y., Cao, J., Ren, H.,& Xu, R. Adaptive calibration: A unified conversion framework of spiking neural networks. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 39, pp. 1583\u20131591 (2025)","DOI":"10.1609\/aaai.v39i2.32150"},{"key":"2872_CR75","doi-asserted-by":"crossref","unstructured":"Wang, Z., Fang, Y., Cao, J., Zhang, Q., Wang, Z., & Xu, R. Masked spiking transformer. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 1761\u20131771 (2023)","DOI":"10.1109\/ICCV51070.2023.00169"},{"key":"2872_CR76","unstructured":"Wang, S., Li, B.Z., Khabsa, M., Fang, H., & Ma, H. Linformer: Self-attention with linear complexity. arXiv:2006.04768 (2020)"},{"key":"2872_CR77","doi-asserted-by":"crossref","unstructured":"Wang, H., Wu, X., Huang, Z., & Xing, E.P. High-frequency component helps explain the generalization of convolutional neural networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8684\u20138694 (2020)","DOI":"10.1109\/CVPR42600.2020.00871"},{"key":"2872_CR78","doi-asserted-by":"crossref","unstructured":"Wang, W., Xie, E., Li, X., Fan, D.-P., Song, K., Liang, D., Lu, T., Luo, P., & Shao, L. Pyramid vision transformer: A versatile backbone for dense prediction without convolutions. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 568\u2013578 (2021)","DOI":"10.1109\/ICCV48922.2021.00061"},{"issue":"10","key":"2872_CR79","doi-asserted-by":"publisher","first-page":"3349","DOI":"10.1109\/TPAMI.2020.2983686","volume":"43","author":"J Wang","year":"2020","unstructured":"Wang, J., Sun, K., Cheng, T., Jiang, B., Deng, C., Zhao, Y., Liu, D., Mu, Y., Tan, M., Wang, X., et al. (2020). Deep high-resolution representation learning for visual recognition. IEEE transactions on pattern analysis and machine intelligence, 43(10), 3349\u20133364.","journal-title":"IEEE transactions on pattern analysis and machine intelligence"},{"key":"2872_CR80","doi-asserted-by":"crossref","unstructured":"Yao, M., Gao, H., Zhao, G., Wang, D., Lin, Y., Yang, Z., & Li, G. Temporal-wise attention spiking neural networks for event streams classification. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 10221\u201310230 (2021)","DOI":"10.1109\/ICCV48922.2021.01006"},{"key":"2872_CR81","unstructured":"Yao, M., Hu, J., Hu, T., Xu, Y., Zhou, Z., Tian, Y., Xu, B., & Li, G. Spike-driven transformer v2: Meta spiking neural network architecture inspiring the design of next-generation neuromorphic chips. arXiv:2404.03663 (2024)"},{"key":"2872_CR82","doi-asserted-by":"crossref","unstructured":"Yao, M., Hu, J., Zhou, Z., Yuan, L., Tian, Y., Xu, B., & Li, G. Spike-driven transformer. Advances in neural information processing systems 36 (2024)","DOI":"10.52202\/075280-2798"},{"key":"2872_CR83","doi-asserted-by":"publisher","first-page":"46","DOI":"10.1016\/j.neucom.2022.11.046","volume":"520","author":"C Ye","year":"2023","unstructured":"Ye, C., Kornijcuk, V., Yoo, D., Kim, J., & Jeong, D. S. (2023). Lacera: Layer-centric event-routing architecture. Neurocomputing, 520, 46\u201359.","journal-title":"Neurocomputing"},{"key":"2872_CR84","doi-asserted-by":"crossref","unstructured":"You, H., Xiong, Y., Dai, X., Wu, B., Zhang, P., Fan, H., Vajda, P., & Lin, Y.C. Castling-vit: Compressing self-attention via switching towards linear-angular attention at vision transformer inference. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 14431\u201314442 (2023)","DOI":"10.1109\/CVPR52729.2023.01387"},{"key":"2872_CR85","doi-asserted-by":"crossref","unstructured":"Yu, W., Luo, M., Zhou, P., Si, C., Zhou, Y., Wang, X., Feng, J., & Yan, S. Metaformer is actually what you need for vision. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 10819\u201310829 (2022)","DOI":"10.1109\/CVPR52688.2022.01055"},{"key":"2872_CR86","doi-asserted-by":"crossref","unstructured":"Zhang, H., Wu, C., Zhang, Z., Zhu, Y., Lin, H., Zhang, Z., Sun, Y., He, T., Mueller, J., Manmatha, R., & et al. Resnest: Split-attention networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 2736\u20132746 (2022)","DOI":"10.1109\/CVPRW56347.2022.00309"},{"key":"2872_CR87","doi-asserted-by":"publisher","first-page":"1229951","DOI":"10.3389\/fnins.2023.1229951","volume":"17","author":"H Zhang","year":"2023","unstructured":"Zhang, H., Li, Y., He, B., Fan, X., Wang, Y., & Zhang, Y. (2023). Direct training high-performance spiking neural networks for object recognition and detection. Frontiers in Neuroscience, 17, 1229951.","journal-title":"Frontiers in Neuroscience"},{"key":"2872_CR88","doi-asserted-by":"crossref","unstructured":"Zheng, H., Wu, Y., Deng, L., Hu, Y., & Li, G. Going deeper with directly-trained larger spiking neural networks. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 35, pp. 11062\u201311070 (2021)","DOI":"10.1609\/aaai.v35i12.17320"},{"key":"2872_CR89","unstructured":"Zhou, Z., Che, K., Fang, W., Tian, K., Zhu, Y., Yan, S., Tian, Y., & Yuan, L. Spikformer v2: Join the high accuracy club on imagenet with an snn ticket. arXiv:2401.02020 (2024)"},{"key":"2872_CR90","unstructured":"Zhou, J., Wang, P., Wang, F., Liu, Q., Li, H., & Jin, R. Elsa: Enhanced local self-attention for vision transformer. arXiv:2112.12786 (2021)"},{"key":"2872_CR91","unstructured":"Zhou, C., Yu, L., Zhou, Z., Ma, Z., Zhang, H., Zhou, H., & Tian, Y. Spikingformer: Spike-driven residual learning for transformer-based spiking neural network. arXiv:2304.11954 (2023)"},{"key":"2872_CR92","doi-asserted-by":"crossref","unstructured":"Zhou, C., Zhang, H., Zhou, Z., Yu, L., Huang, L., Fan, X., Yuan, L., Ma, Z., Zhou, H., & Tian, Y. Qkformer: Hierarchical spiking transformer using qk attention. arXiv:2403.16552 (2024)","DOI":"10.52202\/079017-0416"},{"key":"2872_CR93","doi-asserted-by":"crossref","unstructured":"Zhou, B., Zhao, H., Puig, X., Fidler, S., Barriuso, A., & Torralba, A. Scene parsing through ade20k dataset. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 633\u2013641 (2017)","DOI":"10.1109\/CVPR.2017.544"},{"key":"2872_CR94","unstructured":"Zhou, Z., Zhu, Y., He, C., Wang, Y., Yan, S., Tian, Y., & Yuan, L. Spikformer: When spiking neural network meets transformer. arXiv:2209.15425 (2022)"},{"key":"2872_CR95","unstructured":"Zhu, X., Su, W., Lu, L., Li, B., Wang, X., & Dai, J. Deformable detr: Deformable transformers for end-to-end object detection. arXiv:2010.04159 (2020)"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02872-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-026-02872-6","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-026-02872-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,20]],"date-time":"2026-07-20T16:17:24Z","timestamp":1784564244000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-026-02872-6"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,10]]},"references-count":95,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2872"],"URL":"https:\/\/doi.org\/10.1007\/s11263-026-02872-6","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,10]]},"assertion":[{"value":"11 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 April 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"10 May 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 June 2026","order":5,"name":"change_date","label":"Change Date","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"Update","order":6,"name":"change_type","label":"Change Type","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The original article is revised due to update in affiliation","order":7,"name":"change_details","label":"Change Details","group":{"name":"ArticleHistory","label":"Article History"}}],"article-number":"264"}}