{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,14]],"date-time":"2026-05-14T23:00:29Z","timestamp":1778799629891,"version":"3.51.4"},"reference-count":80,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2025,10,11]],"date-time":"2025-10-11T00:00:00Z","timestamp":1760140800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,10,11]],"date-time":"2025-10-11T00:00:00Z","timestamp":1760140800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100018696","name":"R\u00e9gion Normandie","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100018696","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Machine Vision and Applications"],"published-print":{"date-parts":[[2025,11]]},"DOI":"10.1007\/s00138-025-01748-y","type":"journal-article","created":{"date-parts":[[2025,10,11]],"date-time":"2025-10-11T14:53:29Z","timestamp":1760194409000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Toward simplicity in dynamic inference: a critical study and redesign of early-exit networks"],"prefix":"10.1007","volume":"36","author":[{"given":"Youva","family":"Addad","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Alexis","family":"Lechervy","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Frederic","family":"Jurie","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,10,11]]},"reference":[{"key":"1748_CR1","unstructured":"Tan, M., Le, Q.V.: Efficientnet: rethinking model scaling for convolutional neural networks. In: Chaudhuri, K., Salakhutdinov, R. (eds.) In: Proceedings of the 36th international conference on machine learning, ICML 2019, 9-15 June 2019, Long Beach, California, USA. Proceedings of machine learning research, vol. 97, pp. 6105\u20136114. PMLR, Long Beach, California, USA (2019). http:\/\/proceedings.mlr.press\/v97\/tan19a.html"},{"key":"1748_CR2","unstructured":"Tan, M., Le, Q.V.: Efficientnetv2: Smaller models and faster training. In: Meila, M., Zhang, T. (eds.) In: Proceedings of the 38th international conference on machine learning, ICML 2021, 18-24 July 2021, Virtual Event. Proceedings of machine learning research, vol. 139, pp. 10096\u201310106. PMLR, Online (2021). http:\/\/proceedings.mlr.press\/v139\/tan21a.html"},{"key":"1748_CR3","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (2016)","DOI":"10.1109\/CVPR.2016.90"},{"key":"1748_CR4","doi-asserted-by":"crossref","unstructured":"Huang, G., Liu, Z., Maaten, L., Weinberger, K.Q.: Densely connected convolutional networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (2017)","DOI":"10.1109\/CVPR.2017.243"},{"key":"1748_CR5","doi-asserted-by":"crossref","unstructured":"Han, K., Wang, Y., Tian, Q., Guo, J., Xu, C., Xu, C.: Ghostnet: More features from cheap operations. In: IEEE\/CVF Conference on computer vision and pattern recognition (CVPR) (2020)","DOI":"10.1109\/CVPR42600.2020.00165"},{"key":"1748_CR6","doi-asserted-by":"crossref","unstructured":"Zhang, X., Zhou, X., Lin, M., Sun, J.: Shufflenet: An extremely efficient convolutional neural network for mobile devices. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (2018)","DOI":"10.1109\/CVPR.2018.00716"},{"key":"1748_CR7","unstructured":"Howard, A.G., Zhu, M., Chen, B., Kalenichenko, D., Wang, W., Weyand, T., Andreetto, M., Adam, H.: MobileNets: Efficient convolutional neural networks for mobile vision applications. arXiv. arXiv:1704.04861 [cs] (2017). Accessed 2023-08-29"},{"key":"1748_CR8","doi-asserted-by":"crossref","unstructured":"Howard, A., Sandler, M., Chu, G., Chen, L.-C., Chen, B., Tan, M., Wang, W., Zhu, Y., Pang, R., Vasudevan, V., Le, Q.V., Adam, H.: Searching for mobilenetv3. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (2019)","DOI":"10.1109\/ICCV.2019.00140"},{"key":"1748_CR9","unstructured":"Han, S., Mao, H., Dally, W.J.: Deep compression: Compressing deep neural network with pruning, trained quantization and huffman coding. In: Bengio, Y., LeCun, Y. (eds.) 4th International conference on learning representations, ICLR 2016, San Juan, Puerto Rico, May 2-4, 2016, Conference Track Proceedings (2016). arxiv:1510.00149"},{"key":"1748_CR10","doi-asserted-by":"crossref","unstructured":"He, Y., Liu, P., Wang, Z., Hu, Z., Yang, Y.: Filter pruning via geometric median for deep convolutional neural networks acceleration. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (2019)","DOI":"10.1109\/CVPR42600.2020.00208"},{"key":"1748_CR11","doi-asserted-by":"crossref","unstructured":"Yang, L., Jiang, H., Cai, R., Wang, Y., Song, S., Huang, G., Tian, Q.: Condensenet v2: Sparse feature reactivation for deep networks. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp. 3569\u20133578 (2021)","DOI":"10.1109\/CVPR46437.2021.00357"},{"key":"1748_CR12","doi-asserted-by":"crossref","unstructured":"Liu, Z., Li, J., Shen, Z., Huang, G., Yan, S., Zhang, C.: Learning efficient convolutional networks through network slimming. In: Proceedings of the IEEE International conference on computer vision (ICCV) (2017)","DOI":"10.1109\/ICCV.2017.298"},{"key":"1748_CR13","doi-asserted-by":"crossref","unstructured":"Jung, S., Son, C., Lee, S., Son, J., Han, J.-J., Kwak, Y., Hwang, S.J., Choi, C.: Learning to quantize deep networks by optimizing quantization intervals with task loss. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR) (2019)","DOI":"10.1109\/CVPR.2019.00448"},{"key":"1748_CR14","unstructured":"Hubara, I., Courbariaux, M., Soudry, D., El-Yaniv, R., Bengio, Y.: Binarized neural networks. In: Lee, D.D., Sugiyama, M., Luxburg, U., Guyon, I., Garnett, R. (eds.) Advances in neural information processing systems 29: annual conference on neural information processing systems 2016, December 5-10, 2016, Barcelona, Spain, pp. 4107\u20134115 (2016). https:\/\/proceedings.neurips.cc\/paper\/2016\/hash\/d8330f857a17c53d217014ee776bfd50-Abstract.html"},{"key":"1748_CR15","doi-asserted-by":"publisher","first-page":"525","DOI":"10.1007\/978-3-319-46493-0_32","volume-title":"Comput Vision - ECCV 2016","author":"M Rastegari","year":"2016","unstructured":"Rastegari, M., Ordonez, V., Redmon, J., Farhadi, A.: Xnor-net: Imagenet classification using binary convolutional neural networks. In: Leibe, B., Matas, J., Sebe, N., Welling, M. (eds.) Comput Vision - ECCV 2016, pp. 525\u2013542. Springer, Cham (2016)"},{"key":"1748_CR16","doi-asserted-by":"crossref","unstructured":"Jacob, B., Kligys, S., Chen, B., Zhu, M., Tang, M., Howard, A., Adam, H., Kalenichenko, D.: Quantization and training of neural networks for efficient integer-arithmetic-only inference. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (2018)","DOI":"10.1109\/CVPR.2018.00286"},{"key":"1748_CR17","unstructured":"Hinton, G., Vinyals, O., Dean, J.: Distilling the knowledge in a neural network. In: NIPS deep learning and representation learning workshop (2015). arxiv:1503.02531"},{"key":"1748_CR18","unstructured":"Lan, X., Zhu, X., Gong, S.: Knowledge distillation by on-the-fly native ensemble. In: Bengio, S., Wallach, H.M., Larochelle, H., Grauman, K., Cesa-Bianchi, N., Garnett, R. (eds.) advances in neural information processing systems 31: annual conference on neural information processing systems 2018, NeurIPS 2018, December 3-8, 2018, Montr\u00e9al, Canada, pp. 7528\u20137538 (2018). https:\/\/proceedings.neurips.cc\/paper\/2018\/hash\/94ef7214c4a90790186e255304f8fd1f-Abstract.html"},{"key":"1748_CR19","unstructured":"Ba, J., Caruana, R.: Do deep nets really need to be deep? In: Ghahramani, Z., Welling, M., Cortes, C., Lawrence, N.D., Weinberger, K.Q. (eds.) Advances in neural information processing systems 27: annual conference on neural information processing systems 2014, December 8-13 2014, Montreal, Quebec, Canada, pp. 2654\u20132662 (2014). https:\/\/proceedings.neurips.cc\/paper\/2014\/hash\/ea8fcd92d59581717e06eb187f10666d-Abstract.html"},{"key":"1748_CR20","doi-asserted-by":"crossref","unstructured":"Veit, A., Belongie, S.: Convolutional networks with adaptive inference graphs. In: European conference on computer vision (ECCV) (2018)","DOI":"10.1007\/978-3-030-01246-5_1"},{"key":"1748_CR21","unstructured":"Huang, G., Chen, D., Li, T., Wu, F., Maaten, L., Weinberger, K.: Multi-scale dense networks for resource efficient image classification. In: International conference on learning representations (2018). https:\/\/openreview.net\/forum?id=Hk2aImxAb"},{"key":"1748_CR22","doi-asserted-by":"publisher","unstructured":"Han, Y., Pu, Y., Lai, Z., Wang, C., Song, S., Cao, J., Huang, W., Deng, C., Huang, G.: Learning to weight samples for dynamic early-exiting networks. In: Avidan, S., Brostow, G.J., Ciss\u00e9, M., Farinella, G.M., Hassner, T. (eds.) Computer vision - ECCV 2022 - 17th European Conference, Tel Aviv, Israel, October 23-27, 2022, Proceedings, Part XI. Lecture notes in computer science, vol. 13671, pp. 362\u2013378. Springer, ??? (2022). https:\/\/doi.org\/10.1007\/978-3-031-20083-0_22","DOI":"10.1007\/978-3-031-20083-0_22"},{"key":"1748_CR23","doi-asserted-by":"crossref","unstructured":"Yang, L., Han, Y., Chen, X., Song, S., Dai, J., Huang, G.: Resolution adaptive networks for efficient inference. In: IEEE\/CVF Conference on computer vision and pattern recognition (CVPR) (2020)","DOI":"10.1109\/CVPR42600.2020.00244"},{"key":"1748_CR24","doi-asserted-by":"crossref","unstructured":"Li, H., Zhang, H., Qi, X., Yang, R., Huang, G.: Improved techniques for training adaptive deep networks. In: Proceedings of the IEEE\/CVF International conference on computer vision (ICCV) (2019)","DOI":"10.1109\/ICCV.2019.00198"},{"key":"1748_CR25","doi-asserted-by":"crossref","unstructured":"Phuong, M., Lampert, C.H.: Distillation-based training for multi-exit architectures. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (2019)","DOI":"10.1109\/ICCV.2019.00144"},{"key":"1748_CR26","doi-asserted-by":"crossref","unstructured":"Teerapittayanon, S., McDanel, B., Kung, H.T.: Branchynet: Fast inference via early exiting from deep neural networks. 2016 23rd international conference on pattern recognition (ICPR), 2464\u20132469 (2016)","DOI":"10.1109\/ICPR.2016.7900006"},{"key":"1748_CR27","doi-asserted-by":"publisher","unstructured":"Zhang, L., Song, J., Gao, A., Chen, J., Bao, C., Ma, K.: Be Your Own Teacher: improve the performance of convolutional neural networks via self distillation. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV), pp. 3712\u20133721. IEEE, Seoul, Korea (South) (2019). https:\/\/doi.org\/10.1109\/ICCV.2019.00381","DOI":"10.1109\/ICCV.2019.00381"},{"key":"1748_CR28","doi-asserted-by":"publisher","unstructured":"Lee, H., Lee, J.-S.: Students are the best teacher: exit-ensemble distillation with multi-exits. https:\/\/doi.org\/10.48550\/arXiv.2104.00299 (2021)","DOI":"10.48550\/arXiv.2104.00299"},{"key":"1748_CR29","doi-asserted-by":"crossref","unstructured":"Han, Y., Han, D., Liu, Z., Wang, Y., Pan, X., Pu, Y., Deng, C., Feng, J., Song, S., Huang, G.: Dynamic perceiver for efficient visual recognition (2023)","DOI":"10.1109\/ICCV51070.2023.00551"},{"key":"1748_CR30","unstructured":"Wang, Y., Lv, K., Huang, R., Song, S., Yang, L., Huang, G.: Glance and focus: A dynamic approach to reducing spatial redundancy in image classification. In: Larochelle, H., Ranzato, M., Hadsell, R., Balcan, M.-F., Lin, H.-T. (eds.) Advances in neural information processing systems 33: annual conference on neural information processing systems 2020, NeurIPS 2020, December 6-12, 2020, Virtual (2020)"},{"key":"1748_CR31","unstructured":"Zhu, M., Han, K., Wu, E., Zhang, Q., Nie, Y., Lan, Z., Wang, Y.: Dynamic resolution network. In: Ranzato, M., Beygelzimer, A., Dauphin, Y.N., Liang, P., Vaughan, J.W. (eds.) Advances in neural information processing systems 34: annual conference on neural information processing systems 2021, NeurIPS 2021, December 6-14, 2021, Virtual, pp. 27319\u201327330 (2021). https:\/\/proceedings.neurips.cc\/paper\/2021\/hash\/e56954b4f6347e897f954495eab16a88-Abstract.html"},{"key":"1748_CR32","doi-asserted-by":"publisher","unstructured":"Meronen, L., Trapp, M., Pilzer, A., Yang, L., Solin, A.: Fixing overconfidence in dynamic neural networks. In: IEEE\/CVF Winter conference on applications of computer vision, WACV 2024, Waikoloa, HI, USA, January 3-8, 2024, pp. 2668\u20132678. IEEE, ??? (2024). https:\/\/doi.org\/10.1109\/WACV57701.2024.00266","DOI":"10.1109\/WACV57701.2024.00266"},{"key":"1748_CR33","unstructured":"Wang, Y., Huang, R., Song, S., Huang, Z., Huang, G.: Not all images are worth 16x16 words: Dynamic transformers for efficient image recognition. In: Advances in Neural Information Processing Systems (NeurIPS) (2021)"},{"key":"1748_CR34","doi-asserted-by":"crossref","unstructured":"Chen, M., Lin, M., Li, K., Shen, Y., Wu, Y., Chao, F., Ji, R.: Cf-vit: a general coarse-to-fine method for vision transformer. In: Proceedings of the AAAI conference on artificial intelligence, vol. 37 (2022)","DOI":"10.1609\/aaai.v37i6.25860"},{"key":"1748_CR35","unstructured":"Gao, M.: A survey on recent teacher-student learning studies. arXiv. arXiv:2304.04615 (2023). arxiv:2304.04615 Accessed 2023-08-29"},{"key":"1748_CR36","unstructured":"Sarah, A., Cummings, D., Sridhar, S.N., Sundaresan, S., Szankin, M., Webb, T.J., Mu\u00f1oz, J.P.: A hardware-aware system for accelerating deep neural network optimization. ArXiv abs\/2202.12954 (2022)"},{"key":"1748_CR37","unstructured":"Iandola, F.N., Moskewicz, M.W., Ashraf, K., Han, S., Dally, W.J., Keutzer, K.: Squeezenet: Alexnet-level accuracy with 50x fewer parameters and \u00a11mb model size. ArXiv:1602.07360"},{"issue":"1","key":"1748_CR38","first-page":"1997","volume":"20","author":"T Elsken","year":"2019","unstructured":"Elsken, T., Metzen, J.H., Hutter, F.: Neural architecture search: a survey. J. Mach. Learn. Res. 20(1), 1997\u20132017 (2019)","journal-title":"J. Mach. Learn. Res."},{"key":"1748_CR39","doi-asserted-by":"crossref","unstructured":"Ren, P., Xiao, Y., Chang, X., Huang, P.-Y., Li, Z., Chen, X., Wang, X.: A comprehensive survey of neural architecture search: challenges and solutions. arXiv. arXiv:2006.02903 [cs, stat] (2021)","DOI":"10.1145\/3447582"},{"issue":"11","key":"1748_CR40","doi-asserted-by":"publisher","first-page":"7436","DOI":"10.1109\/TPAMI.2021.3117837","volume":"44","author":"Y Han","year":"2022","unstructured":"Han, Y., Huang, G., Song, S., Yang, L., Wang, H., Wang, Y.: Dynamic neural networks: a survey. IEEE Trans Pattern Anal Mach Intell 44(11), 7436\u20137456 (2022). https:\/\/doi.org\/10.1109\/TPAMI.2021.3117837","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"1748_CR41","doi-asserted-by":"publisher","unstructured":"Laskaridis, S., Kouris, A., Lane, N.D.: Adaptive inference through early-exit networks: Design, challenges and directions. In: Proceedings of the 5th international workshop on embedded and mobile deep learning. EMDL\u201921, pp. 1\u20136. Association for computing machinery, New York, NY, USA (2021). https:\/\/doi.org\/10.1145\/3469116.3470012","DOI":"10.1145\/3469116.3470012"},{"issue":"1","key":"1748_CR42","doi-asserted-by":"publisher","first-page":"5","DOI":"10.1109\/TEVC.2022.3220747","volume":"27","author":"Y Bi","year":"2023","unstructured":"Bi, Y., Xue, B., Mesejo, P., Cagnoni, S., Zhang, M.: A survey on evolutionary computation for computer vision and image analysis: past, present, and future trends. IEEE Trans Evolut Comput 27(1), 5\u201325 (2023). https:\/\/doi.org\/10.1109\/TEVC.2022.3220747. arXiv:2209.06399 [cs].","journal-title":"IEEE Trans Evolut Comput"},{"key":"1748_CR43","doi-asserted-by":"publisher","unstructured":"Hedderich, M.A., Lange, L., Adel, H., Str\u00f6tgen, J., Klakow, D.: A survey on recent approaches for natural language processing in low-resource scenarios. In: Proceedings of the 2021 conference of the North American chapter of the association for computational linguistics: human language technologies, pp. 2545\u20132568. Association for Computational Linguistics, Online (2021). https:\/\/doi.org\/10.18653\/v1\/2021.naacl-main.201 . https:\/\/aclanthology.org\/2021.naacl-main.201","DOI":"10.18653\/v1\/2021.naacl-main.201"},{"key":"1748_CR44","unstructured":"Prabhavalkar, R., Hori, T., Sainath, T.N., Schl\u00fcter, R., Watanabe, S.: End-to-end speech recognition: a survey. arXiv. arXiv:2303.03329 [cs, eess] (2023). Accessed 2023-08-29"},{"key":"1748_CR45","doi-asserted-by":"publisher","unstructured":"Yu, H., Li, H., Hua, G., Huang, G., Shi, H.: Boosted dynamic neural networks. In: Williams, B., Chen, Y., Neville, J. (eds.) In: Thirty-Seventh AAAI conference on artificial intelligence, AAAI 2023, thirty-fifth conference on innovative applications of artificial intelligence, IAAI 2023, thirteenth symposium on educational advances in artificial intelligence, EAAI 2023, Washington, DC, USA, February 7-14, 2023, pp. 10989\u201310997. AAAI Press, Washington, DC, USA (2023). https:\/\/doi.org\/10.1609\/AAAI.V37I9.26302","DOI":"10.1609\/AAAI.V37I9.26302"},{"key":"1748_CR46","doi-asserted-by":"crossref","unstructured":"Addad, Y., Lechervy, A., Jurie, F.: Multi-exit resource-efficient neural architecture for image classification with optimized fusion block. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) Workshops, pp. 1486\u20131491 (2023)","DOI":"10.1109\/ICCVW60793.2023.00161"},{"key":"1748_CR47","doi-asserted-by":"publisher","unstructured":"Agarwal, C., D\u2019souza, D., Hooker, S.: Estimating example difficulty using variance of gradients. In: 2022 IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp. 10358\u201310368. IEEE, New Orleans, LA, USA (2022). https:\/\/doi.org\/10.1109\/CVPR52688.2022.01012 . https:\/\/ieeexplore.ieee.org\/document\/9879352\/ Accessed 2023-03-02","DOI":"10.1109\/CVPR52688.2022.01012"},{"key":"1748_CR48","doi-asserted-by":"publisher","unstructured":"Lee, H., Lee, J.: Rethinking online knowledge distillation with multi-exits. In: Wang, L., Gall, J., Chin, T., Sato, I., Chellappa, R. (eds.) Computer Vision - ACCV 2022 - 16th Asian conference on computer vision, Macao, China, December 4-8, 2022, Proceedings, Part VI. Lecture notes in computer science, vol. 13846, pp. 408\u2013424. Springer, Macao, China (2022). https:\/\/doi.org\/10.1007\/978-3-031-26351-4_25","DOI":"10.1007\/978-3-031-26351-4_25"},{"key":"1748_CR49","doi-asserted-by":"publisher","unstructured":"Gong, C., Chen, Y., Luo, Q., Lu, Y., Li, T., Zhang, Y., Sun, Y., Zhang, L.: Deep feature surgery: Towards accurate and efficient multi-exit networks. In: Leonardis, A., Ricci, E., Roth, S., Russakovsky, O., Sattler, T., Varol, G. (eds.) Computer Vision - ECCV 2024 - 18th European Conference, Milan, Italy, September 29-October 4, 2024, Proceedings, Part XLIX. Lecture Notes in Computer Science, vol. 15107, pp. 435\u2013451. Springer, Milan, Italy (2024). https:\/\/doi.org\/10.1007\/978-3-031-72967-6_24","DOI":"10.1007\/978-3-031-72967-6_24"},{"key":"1748_CR50","doi-asserted-by":"publisher","first-page":"17","DOI":"10.1007\/978-3-031-78172-8_2","volume-title":"Pattern Recognition","author":"Y Addad","year":"2025","unstructured":"Addad, Y., Lechervy, A., Jurie, F.: Balancing accuracy and efficiency in budget-aware early-exiting neural networks. In: Antonacopoulos, A., Chaudhuri, S., Chellappa, R., Liu, C.-L., Bhattacharya, S., Pal, U. (eds.) Pattern Recognition, pp. 17\u201331. Springer, Cham (2025)"},{"key":"1748_CR51","doi-asserted-by":"publisher","unstructured":"Ilhan, F., Chow, K., Hu, S., Huang, T., Tekin, S.F., Wei, W., Wu, Y., Lee, M., Kompella, R., Latapie, H., Liu, G., Liu, L.: Adaptive deep neural network inference optimization with eenet. In: IEEE\/CVF Winter conference on applications of computer vision, WACV 2024, Waikoloa, HI, USA, January 3-8, 2024, pp. 1362\u20131371. IEEE, ??? (2024). https:\/\/doi.org\/10.1109\/WACV57701.2024.00140","DOI":"10.1109\/WACV57701.2024.00140"},{"key":"1748_CR52","doi-asserted-by":"crossref","unstructured":"Addad, Y., Lechervy, A., Jurie, F.: CHASE: Channel-wise and spatial attention for early exiting in image classification. Working paper or preprint (2025). https:\/\/hal.science\/hal-04885744","DOI":"10.1109\/ICASSP49660.2025.10888714"},{"key":"1748_CR53","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, L., Polosukhin, I.: Attention is all you need. In: Guyon, I., Luxburg, U., Bengio, S., Wallach, H.M., Fergus, R., Vishwanathan, S.V.N., Garnett, R. (eds.) In: Advances in neural information processing systems 30: annual conference on neural information processing systems 2017, December 4-9, 2017, Long Beach, CA, USA, Long Beach, CA, USA, pp. 5998\u20136008 (2017). https:\/\/proceedings.neurips.cc\/paper\/2017\/hash\/3f5ee243547dee91fbd053c1c4a845aa-Abstract.html"},{"key":"1748_CR54","doi-asserted-by":"publisher","unstructured":"Liu, Z., Mao, H., Wu, C., Feichtenhofer, C., Darrell, T., Xie, S.: A convnet for the 2020s. In: IEEE\/CVF Conference on computer vision and pattern recognition, CVPR 2022, New Orleans, LA, USA, June 18-24, 2022, pp. 11966\u201311976. IEEE, New Orleans, LA, USA (2022). https:\/\/doi.org\/10.1109\/CVPR52688.2022.01167","DOI":"10.1109\/CVPR52688.2022.01167"},{"key":"1748_CR55","unstructured":"Vasu, P.K.A., Gabriel, J., Zhu, J., Tuzel, O., Ranjan, A.: Fastvit: A fast hybrid vision transformer using structural reparameterization. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV), pp. 5785\u20135795 (2023)"},{"issue":"2","key":"1748_CR56","doi-asserted-by":"publisher","first-page":"896","DOI":"10.1109\/TPAMI.2023.3329173","volume":"46","author":"W Yu","year":"2024","unstructured":"Yu, W., Si, C., Zhou, P., Luo, M., Zhou, Y., Feng, J., Yan, S., Wang, X.: Metaformer baselines for vision. IEEE Trans. Pattern Anal. Mach. Intell. 46(2), 896\u2013912 (2024). https:\/\/doi.org\/10.1109\/TPAMI.2023.3329173","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"1748_CR57","unstructured":"Chu, X., Tian, Z., Wang, Y., Zhang, B., Ren, H., Wei, X., Xia, H., Shen, C.: Twins: Revisiting the design of spatial attention in vision transformers. In: Ranzato, M., Beygelzimer, A., Dauphin, Y.N., Liang, P., Vaughan, J.W. (eds.) In: Advances in neural information processing systems 34: annual conference on neural information processing systems 2021, NeurIPS 2021, December 6-14, 2021, Virtual, pp. 9355\u20139366 (2021). https:\/\/proceedings.neurips.cc\/paper\/2021\/hash\/4e0928de075538c593fbdabb0c5ef2c3-Abstract.html"},{"key":"1748_CR58","unstructured":"Chu, X., Tian, Z., Zhang, B., Wang, X., Shen, C.: Conditional positional encodings for vision transformers. In: The eleventh international conference on learning representations, ICLR 2023, Kigali, Rwanda, May 1-5, 2023. OpenReview.net, Kigali, Rwanda (2023). https:\/\/openreview.net\/pdf?id=3KWnuT-R1bh"},{"key":"1748_CR59","doi-asserted-by":"crossref","unstructured":"Hu, J., Shen, L., Sun, G.: Squeeze-and-excitation networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (2018)","DOI":"10.1109\/CVPR.2018.00745"},{"key":"1748_CR60","unstructured":"Krizhevsky, A.: Learning multiple layers of features from tiny images, 32\u201333 (2009)"},{"key":"1748_CR61","doi-asserted-by":"crossref","unstructured":"Deng, J., Dong, W., Socher, R., Li, L.-J., Li, K., Fei-Fei, L.: Imagenet: A large-scale hierarchical image database. In: 2009 IEEE conference on computer vision and pattern recognition, pp. 248\u2013255 (2009). IEEE","DOI":"10.1109\/CVPR.2009.5206848"},{"key":"1748_CR62","unstructured":"Bao, H., Dong, L., Piao, S., Wei, F.: BEit: BERT pre-training of image transformers. In: International conference on learning representations (2022). https:\/\/openreview.net\/forum?id=p-BhZSz59o4"},{"key":"1748_CR63","unstructured":"Cubuk, E.D., Zoph, B., Shlens, J., Le, Q.: Randaugment: Practical automated data augmentation with a reduced search space. In: Larochelle, H., Ranzato, M., Hadsell, R., Balcan, M., Lin, H. (eds.) In: Advances in neural information processing systems 33: annual conference on neural information processing systems 2020, NeurIPS 2020, December 6-12, 2020, Virtual, Online (2020). https:\/\/proceedings.neurips.cc\/paper\/2020\/hash\/d85b63ef0ccb114d0a3bb7b7d808028f-Abstract.html"},{"key":"1748_CR64","doi-asserted-by":"crossref","unstructured":"Szegedy, C., Vanhoucke, V., Ioffe, S., Shlens, J., Wojna, Z.: Rethinking the inception architecture for computer vision. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR) (2016)","DOI":"10.1109\/CVPR.2016.308"},{"key":"1748_CR65","unstructured":"Zhang, H., Cisse, M., Dauphin, Y.N., Lopez-Paz, D.: mixup: Beyond empirical risk minimization. In: International conference on learning representations (2018). https:\/\/openreview.net\/forum?id=r1Ddp1-Rb"},{"key":"1748_CR66","doi-asserted-by":"crossref","unstructured":"Yun, S., Han, D., Oh, S.J., Chun, S., Choe, J., Yoo, Y.: Cutmix: Regularization strategy to train strong classifiers with localizable features. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (2019)","DOI":"10.1109\/ICCV.2019.00612"},{"key":"1748_CR67","doi-asserted-by":"publisher","first-page":"646","DOI":"10.1007\/978-3-319-46493-0_39","volume-title":"Computer Vision - ECCV 2016","author":"G Huang","year":"2016","unstructured":"Huang, G., Sun, Y., Liu, Z., Sedra, D., Weinberger, K.Q.: Deep networks with stochastic depth. In: Leibe, B., Matas, J., Sebe, N., Welling, M. (eds.) Computer Vision - ECCV 2016, pp. 646\u2013661. Springer, Cham (2016)"},{"key":"1748_CR68","doi-asserted-by":"crossref","unstructured":"Touvron, H., Cord, M., Sablayrolles, A., Synnaeve, G., J\u00e9gou, H.: Going deeper with image transformers. In: Proceedings of the IEEE\/CVF International conference on computer vision (ICCV), pp. 32\u201342 (2021)","DOI":"10.1109\/ICCV48922.2021.00010"},{"issue":"4","key":"1748_CR69","doi-asserted-by":"publisher","first-page":"838","DOI":"10.1137\/0330046","volume":"30","author":"BT Polyak","year":"1992","unstructured":"Polyak, B.T., Juditsky, A.B.: Acceleration of stochastic approximation by averaging. SIAM J Control Optim 30(4), 838\u2013855 (1992). https:\/\/doi.org\/10.1137\/0330046","journal-title":"SIAM J Control Optim"},{"key":"1748_CR70","unstructured":"Zhong, Z., Zheng, L., Kang, G., Li, S., Yang, Y.: Random erasing data augmentation. arXiv:1708.04896 (2017)"},{"key":"1748_CR71","doi-asserted-by":"crossref","unstructured":"Radosavovic, I., Kosaraju, R.P., Girshick, R., He, K., Dollar, P.: Designing network design spaces. In: IEEE\/CVF Conference on computer vision and pattern recognition (CVPR) (2020)","DOI":"10.1109\/CVPR42600.2020.01044"},{"key":"1748_CR72","unstructured":"Touvron, H., Cord, M., Douze, M., Massa, F., Sablayrolles, A., J\u00e9gou, H.: Training data-efficient image transformers & distillation through attention. In: Meila, M., Zhang, T. (eds.) In: Proceedings of the 38th international conference on machine learning, ICML 2021, 18-24 July 2021, Virtual event. proceedings of machine learning research, vol. 139, pp. 10347\u201310357. PMLR, Online (2021). http:\/\/proceedings.mlr.press\/v139\/touvron21a.html"},{"key":"1748_CR73","doi-asserted-by":"crossref","unstructured":"Yu, W., Luo, M., Zhou, P., Si, C., Zhou, Y., Wang, X., Feng, J., Yan, S.: Metaformer is actually what you need for vision. In: Proceedings of the IEEE\/CVF Conference on computer vision and pattern recognition (CVPR), pp. 10819\u201310829 (2022)","DOI":"10.1109\/CVPR52688.2022.01055"},{"key":"1748_CR74","unstructured":"Mehta, S., Rastegari, M.: Mobilevit: Light-weight, general-purpose, and mobile-friendly vision transformer. In: International conference on learning representations (2022). https:\/\/openreview.net\/forum?id=vh-0sUt8HlG"},{"key":"1748_CR75","doi-asserted-by":"crossref","unstructured":"Pan, J., Bulat, A., Tan, F., Zhu, X., Dudziak, L., Li, H., Tzimiropoulos, G., Martinez, B.: Edgevits: Competing light-weight cnns on mobile devices with vision transformers. In: European conference on computer vision (2022)","DOI":"10.1007\/978-3-031-20083-0_18"},{"key":"1748_CR76","first-page":"12934","volume":"35","author":"Y Li","year":"2022","unstructured":"Li, Y., Yuan, G., Wen, Y., Hu, J., Evangelidis, G., Tulyakov, S., Wang, Y., Ren, J.: Efficientformer: vision transformers at mobilenet speed. Adv. Neural. Inf. Process. Syst. 35, 12934\u201312949 (2022)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"1748_CR77","doi-asserted-by":"crossref","unstructured":"Chen, J., Kao, S.-h., He, H., Zhuo, W., Wen, S., Lee, C.-H., Chan, S.-H.G.: Run, don\u2019t walk: Chasing higher flops for faster neural networks. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp. 12021\u201312031 (2023)","DOI":"10.1109\/CVPR52729.2023.01157"},{"key":"1748_CR78","doi-asserted-by":"crossref","unstructured":"Vasu, P.K.A., Gabriel, J., Zhu, J., Tuzel, O., Ranjan, A.: Mobileone: An improved one millisecond mobile backbone. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp. 7907\u20137917 (2023)","DOI":"10.1109\/CVPR52729.2023.00764"},{"key":"1748_CR79","doi-asserted-by":"crossref","unstructured":"Wang, J., Zhang, S., Liu, Y., Wu, T., Yang, Y., Liu, X., Chen, K., Luo, P., Lin, D.: Riformer: Keep your vision backbone effective but removing token mixer. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR), pp. 14443\u201314452 (2023)","DOI":"10.1109\/CVPR52729.2023.01388"},{"key":"1748_CR80","doi-asserted-by":"crossref","unstructured":"Shaker, A., Maaz, M., Rasheed, H., Khan, S., Yang, M.-H., Khan, F.S.: Swiftformer: Efficient additive attention for transformer-based real-time mobile vision applications. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV) (2023)","DOI":"10.1109\/ICCV51070.2023.01598"}],"container-title":["Machine Vision and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00138-025-01748-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00138-025-01748-y\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00138-025-01748-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,11,10]],"date-time":"2025-11-10T17:01:33Z","timestamp":1762794093000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00138-025-01748-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,10,11]]},"references-count":80,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2025,11]]}},"alternative-id":["1748"],"URL":"https:\/\/doi.org\/10.1007\/s00138-025-01748-y","relation":{"has-preprint":[{"id-type":"doi","id":"10.21203\/rs.3.rs-6608894\/v1","asserted-by":"object"}]},"ISSN":["0932-8092","1432-1769"],"issn-type":[{"value":"0932-8092","type":"print"},{"value":"1432-1769","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,10,11]]},"assertion":[{"value":"7 May 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 August 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"22 September 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"11 October 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"127"}}