{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,12]],"date-time":"2026-06-12T19:09:10Z","timestamp":1781291350883,"version":"3.54.1"},"reference-count":46,"publisher":"Springer Science and Business Media LLC","issue":"10","license":[{"start":{"date-parts":[[2024,10,1]],"date-time":"2024-10-01T00:00:00Z","timestamp":1727740800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,10,1]],"date-time":"2024-10-01T00:00:00Z","timestamp":1727740800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["J Sign Process Syst"],"published-print":{"date-parts":[[2024,10]]},"DOI":"10.1007\/s11265-024-01929-4","type":"journal-article","created":{"date-parts":[[2024,11,1]],"date-time":"2024-11-01T10:03:42Z","timestamp":1730455422000},"page":"569-585","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Age Estimation from Speech Using Tuned CNN Model on Edge Devices"],"prefix":"10.1007","volume":"96","author":[{"given":"Laxmi Kantham","family":"Durgam","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Ravi Kumar","family":"Jatoth","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,11,1]]},"reference":[{"key":"1929_CR1","doi-asserted-by":"publisher","first-page":"3535","DOI":"10.1007\/s11042-021-11614-4","volume":"81","author":"HA S\u00e1nchez-Hevia","year":"2022","unstructured":"S\u00e1nchez-Hevia, H. A., Gil-Pita, R., Utrilla-Manso, M., & Rosa-Zurera, M. (2022). Age group classification and gender recognition from speech with temporal convolutional neural networks. Multimedia Tools and Applications, 81, 3535\u20133552.","journal-title":"Multimedia Tools and Applications"},{"key":"1929_CR2","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1049\/sil2.12216","volume":"17","author":"CH Lin","year":"2023","unstructured":"Lin, C. H., Lai, H. Y., Huang, P. T., Chen, P. Y., & Li, C. M. (2023). Vowel classification with combining pitch detection and one-dimensional convolutional neural network based classifier for gender identification. IET Signal Processing, 17, 1\u201314. https:\/\/doi.org\/10.1049\/sil2.12216","journal-title":"IET Signal Processing"},{"key":"1929_CR3","doi-asserted-by":"publisher","first-page":"5655","DOI":"10.1007\/s12652-021-03238-1","volume":"13","author":"K Kuppusamy","year":"2021","unstructured":"Kuppusamy, K., & Eswaran, C. (2021). Convolutional and deep neural networks based techniques for extracting the age-relevant features of the speaker. Journal of Ambient Intelligence and Humanized Computing, 13, 5655\u20135667.","journal-title":"Journal of Ambient Intelligence and Humanized Computing"},{"issue":"10","key":"1929_CR4","first-page":"1533","volume":"22","author":"O Abdel-Hamid","year":"2014","unstructured":"Abdel-Hamid, O., Mohamed, A. R., Jiang, H., Deng, L., Penn, G., & Yu, D. (2014). Convolutional neural networks for speech recognition. IEEE ACM Transactions on audio, speech, and language processing, 22(10), 1533\u20131545.","journal-title":"IEEE ACM Transactions on audio, speech, and language processing"},{"key":"1929_CR5","doi-asserted-by":"publisher","first-page":"50285","DOI":"10.1109\/ACCESS.2023.3278106","volume":"11","author":"Z Hu","year":"2023","unstructured":"Hu, Z., LingHu, K., Yu, H., & Liao, C. (2023). Speech emotion recognition based on attention mcnn combined with gender information. IEEE Access, 11, 50285\u201350294.","journal-title":"IEEE Access"},{"key":"1929_CR6","doi-asserted-by":"publisher","first-page":"2225333","DOI":"10.1109\/ACCESS.2020.3043894","volume":"8","author":"S Zhong","year":"2020","unstructured":"Zhong, S., Yu, B., & Zhang, H. (2020). Exploration of an independent training framework for speech emotion recognition. IEEE Access, 8, 2225333\u20132225343.","journal-title":"IEEE Access"},{"issue":"5","key":"1929_CR7","doi-asserted-by":"publisher","first-page":"991","DOI":"10.1109\/TCYB.2014.2341737","volume":"45","author":"J Yu","year":"2015","unstructured":"Yu, J., & Wang, Z.-F. (2015). A video, text, and speech-driven realistic 3-d virtual head for human\u2013machine interface. IEEE Transactions on Cybernetics, 45(5), 991\u20131002.","journal-title":"IEEE Transactions on Cybernetics"},{"key":"1929_CR8","doi-asserted-by":"publisher","first-page":"79236","DOI":"10.1109\/ACCESS.2021.3084299","volume":"9","author":"MM Kabir","year":"2021","unstructured":"Kabir, M. M., Mridha, M. F., Shin, J., Jahan, I., & Ohi, A. Q. (2021). A survey of speaker recognition: Fundamental theories, recognition methods and opportunities. IEEE Access, 9, 79236\u201379263.","journal-title":"IEEE Access"},{"key":"1929_CR9","doi-asserted-by":"publisher","first-page":"182348","DOI":"10.1109\/ACCESS.2019.2959906","volume":"7","author":"H Huang","year":"2019","unstructured":"Huang, H., Wang, X., Hu, M., & Tao, Y. (2019). Applied to mobile multimedia intelligent speech system interactive topic guiding model. IEEE Access, 7, 182348\u2013182356.","journal-title":"IEEE Access"},{"key":"1929_CR10","doi-asserted-by":"publisher","first-page":"441","DOI":"10.1016\/j.csl.2019.06.001","volume":"58","author":"A Nautsch","year":"2019","unstructured":"Nautsch, A., Jim\u00e9nez, A., Treiber, A., Kolberg, J., Jasserand, C., Kindt, E., Delgado, H., Todisco, M., & Hmani, M. (2019). Abdelraheem: Preserving privacy in speaker and speech characterisation. Computer Speech & Language, 58, 441\u2013480.","journal-title":"Computer Speech & Language"},{"key":"1929_CR11","doi-asserted-by":"publisher","first-page":"367","DOI":"10.1007\/s10772-021-09808-0","volume":"24","author":"KB Bhangale","year":"2021","unstructured":"Bhangale, K. B., & Mohanaprasad, K. (2021). A review on speech processing using machine learning paradigm. International Journal of Speech Technology, 24, 367\u2013388.","journal-title":"International Journal of Speech Technology"},{"issue":"4","key":"1929_CR12","doi-asserted-by":"publisher","first-page":"1595","DOI":"10.1016\/j.jksuci.2021.11.019","volume":"34","author":"KB Bhangale","year":"2022","unstructured":"Bhangale, K. B., & Mohanaprasad, K. (2022). A review on tinyml: State-of-the-art and prospects. Journal of King Saud University-Computer and Information Sciences, 34(4), 1595\u20131623.","journal-title":"Journal of King Saud University-Computer and Information Sciences"},{"key":"1929_CR13","unstructured":"Warden, P., & Situnayake, D. (2019). Tinyml: Machine learning with tensorflow lite on arduino and ultra-low-power microcontrollers. O\u2019Reilly Media, 1595\u20131623."},{"key":"1929_CR14","doi-asserted-by":"publisher","first-page":"73484","DOI":"10.1109\/ACCESS.2022.3189776","volume":"10","author":"E Manor","year":"2022","unstructured":"Manor, E., & Greenberg, S. (2022). Custom hardware inference accelerator fortensorflow lite for microcontrollers. IEEE Access, 10, 73484\u201373484.","journal-title":"IEEE Access"},{"key":"1929_CR15","doi-asserted-by":"publisher","first-page":"5","DOI":"10.1016\/j.knosys.2016.10.008","volume":"115","author":"Z Qawaqneh","year":"2017","unstructured":"Qawaqneh, Z., Mallouh, A. A., & Barkana, B. D. (2017). Deep neural network framework and transformed mfccs for speaker\u2019s age and gender classification. Knowledge-Based Systems, 115, 5\u201314.","journal-title":"Knowledge-Based Systems"},{"key":"1929_CR16","doi-asserted-by":"publisher","first-page":"329","DOI":"10.1007\/s12652-019-01303-4","volume":"11","author":"GK Birajdar","year":"2020","unstructured":"Birajdar, G. K., & Patil, M. D. (2020). Speech music classification using visual and spectral chromagram features. Journal of Ambient Intelligence and Humanized Computing, 11, 329\u2013347.","journal-title":"Journal of Ambient Intelligence and Humanized Computing"},{"issue":"2","key":"1929_CR17","doi-asserted-by":"publisher","first-page":"130","DOI":"10.1109\/LSP.2010.2100380","volume":"18","author":"J Dennis","year":"2011","unstructured":"Dennis, J., Tran, H. D., & Li, H. (2011). Spectrogram image feature for sound event classification in mismatched conditions. IEEE Signal Processing Letters, 18(2), 130\u2013133.","journal-title":"IEEE Signal Processing Letters"},{"key":"1929_CR18","unstructured":"impulse: https:\/\/docs.edgeimpulse.com\/docs"},{"key":"1929_CR19","unstructured":"Bagur, J. (2023). Edge Impulse with the Nano 33 BLE Sense. https:\/\/docs.arduino.cc\/tutorials\/nano-33-ble-sense\/edge-impulse"},{"key":"1929_CR20","unstructured":"Nano, J.: https:\/\/developer.nvidia.com\/embedded\/jetson-modules"},{"key":"1929_CR21","doi-asserted-by":"publisher","first-page":"211","DOI":"10.1016\/j.icte.2021.12.007","volume":"9","author":"KH Le","year":"2021","unstructured":"Le, K. H., Le-Minh, K. H., & Thai, H. T. (2021). Brainyedge: An ai-enabled framework for iot edge computing. ICT Express, 9, 211\u2013221.","journal-title":"ICT Express"},{"issue":"11","key":"1929_CR22","doi-asserted-by":"publisher","first-page":"3897","DOI":"10.1109\/TCSI.2018.2852260","volume":"65","author":"A Ibrahim","year":"2018","unstructured":"Ibrahim, A., & Valle, M. (2018). Real-time embedded machine learning for tensorial tactile data processing. IEEE Transactions on Circuits and Systems I: Regular Papers, 65(11), 3897\u20133906. https:\/\/doi.org\/10.1109\/TCSI.2018.2852260","journal-title":"IEEE Transactions on Circuits and Systems I: Regular Papers"},{"key":"1929_CR23","doi-asserted-by":"crossref","unstructured":"Maayah, M., Abunada, A., Al-Janahi, K., Ahmed, M. E., & Qadir, J. (2023). Limitaccess: on-device tinyml based robust speech recognition and age classification. Discover Artificial Intelligence, 3(8).","DOI":"10.1007\/s44163-023-00051-x"},{"key":"1929_CR24","doi-asserted-by":"crossref","unstructured":"Kennedy, J., Lemaignan, S., Montassier, C., Lavalade, P., Irfan, B., Papadopoulos, F., Senft, E. & Belpaeme, T. (2016). Children speech recording data set, human-robot interaction.","DOI":"10.1145\/2909824.3020229"},{"key":"1929_CR25","doi-asserted-by":"publisher","unstructured":"Iloanusi, O., Ejiogu, U., Okoye, I.E., Ezika, I., Ezichi, S., Osuagwu, C., & Ejiogu, E. (2019). Voice Recognition and Gender Classification in the Context of Native Languages and Lingua Franca,. Paper presented at the 6th International Conference on Soft Computing & Machine Intelligence , Johannesburg, South Africa. https:\/\/doi.org\/10.1109\/ISCMI47871.2019.9004306.","DOI":"10.1109\/ISCMI47871.2019.9004306"},{"key":"1929_CR26","unstructured":"m4, A.: https:\/\/www.arm.com\/products\/silicon-ip-cpu\/cortex-m\/cortex-m4"},{"key":"1929_CR27","doi-asserted-by":"crossref","unstructured":"Mao, D., Sun, H., Li, X., Yu, X., Wu, J., & Zhang, Q. (2023). Real-time fruit detection using deep neural networks on cpu (rtfd): An edge ai application. Computers and Electronics in Agriculture, 204.","DOI":"10.1016\/j.compag.2022.107517"},{"key":"1929_CR28","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1016\/j.iot.2022.100670","volume":"21","author":"N Tekin","year":"2023","unstructured":"Tekin, N., Acar, A., Aris, A., Uluagac, A. S., & Gungor, V. C. (2023). Energy consumption of on-device machine learning models for iot intrusion detection. Internet of Things, 21, 1.","journal-title":"Internet of Things"},{"issue":"8","key":"1929_CR29","doi-asserted-by":"publisher","first-page":"11897","DOI":"10.1007\/s11042-022-13725-y","volume":"82","author":"S Patnaik","year":"2023","unstructured":"Patnaik, S. (2023). Speech emotion recognition by using complex mfcc and deep sequential model. Multimedia Tools and Applications., 82(8), 11897\u20131192.","journal-title":"Multimedia Tools and Applications."},{"key":"1929_CR30","doi-asserted-by":"publisher","first-page":"37","DOI":"10.1016\/j.procs.2019.04.009","volume":"151","author":"FA Shaqra","year":"2019","unstructured":"Shaqra, F. A., Duwairi, R., & Al-Ayyoub, M. (2019). Recognizing emotion from speech based on age and gender using hierarchical models. Procedia Computer Science, 151, 37\u201344.","journal-title":"Procedia Computer Science"},{"key":"1929_CR31","doi-asserted-by":"crossref","unstructured":"Kang, Z., Wang, J., Peng, J., & Xiao, J. (2023). Svldl: Improved speaker age estimation using selective variance label distribution learning. IEEE Spoken Language Technology Workshop (SLT), 1037\u20131044.","DOI":"10.1109\/SLT54892.2023.10023124"},{"key":"1929_CR32","doi-asserted-by":"publisher","first-page":"312","DOI":"10.1016\/j.bspc.2018.08.035","volume":"47","author":"J Zhao","year":"2019","unstructured":"Zhao, J., Mao, X., & Chen, L. (2019). Speech emotion recognition using deep 1d & 2d cnn lstm networks. Biomedical Signal Processing and Control, 47, 312\u2013323.","journal-title":"Biomedical Signal Processing and Control"},{"key":"1929_CR33","doi-asserted-by":"publisher","first-page":"13951","DOI":"10.1007\/s00521-022-07246-w","volume":"34","author":"M Subramanian","year":"2022","unstructured":"Subramanian, M., Shanmugavadivel, K., & Nandhini, P. S. (2022). On fine-tuning deep learning models using transfer learning and hyper-parameters optimization for disease identification in maize leaves. Neural Computing and Applications, 34, 13951\u201313968.","journal-title":"Neural Computing and Applications"},{"key":"1929_CR34","doi-asserted-by":"publisher","first-page":"316","DOI":"10.1016\/j.procs.2017.08.003","volume":"112","author":"I Rebai","year":"2017","unstructured":"Rebai, I., BenAyed, Y., Mahdi, W., & Lorr\u00e9, J. P. (2017). Improving speech recognition using data augmentation and acoustic model fusion. Procedia Computer Science, 112, 316\u2013322.","journal-title":"Procedia Computer Science"},{"key":"1929_CR35","unstructured":"Salman, S., Liu, X. (2019). Overfitting mechanism and avoidance in deep neural networks. arXiv:1901.06566"},{"key":"1929_CR36","doi-asserted-by":"publisher","first-page":"107647","DOI":"10.1016\/j.apacoust.2020.107647","volume":"172","author":"Z Wang","year":"2021","unstructured":"Wang, Z., Zhang, T., Shao, Y., & Ding, B. (2021). Lstm convolutional blstm encoder decoder network for minimum mean-square error approach to speech enhancement. Applied Acoustics, 172, 107647.","journal-title":"Applied Acoustics"},{"key":"1929_CR37","doi-asserted-by":"publisher","first-page":"216","DOI":"10.1016\/j.patcog.2019.02.023","volume":"91","author":"A Luque","year":"2019","unstructured":"Luque, A., Carrasco, A., Mart\u00edn, A., & de Las, Heras A. (2019). The impact of class imbalance in classification performance metrics based on the binary confusion matrix. Pattern Recognition, 91, 216\u2013231.","journal-title":"Pattern Recognition"},{"key":"1929_CR38","doi-asserted-by":"crossref","unstructured":"Kang, W., & Chung, J. (2019). Power- and time-aware deep learning inference for mobile embedded devices. Pattern Recognition, 7, 3778\u20133789.","DOI":"10.1109\/ACCESS.2018.2887099"},{"key":"1929_CR39","doi-asserted-by":"publisher","unstructured":"Koc, W. W., Chang, Y. T., Yu, J. Y., & \u0130k, T. U. (2021). Text-to-speech with model compression on edge devices., 25, 114\u2013119. https:\/\/doi.org\/10.23919\/APNOMS52696.2021.9562651. 22nd Asia-Pacific Network Operations and Management Symposium (APNOMS), Tainan.","DOI":"10.23919\/APNOMS52696.2021.9562651"},{"key":"1929_CR40","doi-asserted-by":"publisher","first-page":"268","DOI":"10.1016\/j.csl.2017.06.002","volume":"46","author":"H Kaya","year":"2017","unstructured":"Kaya, H., Salah, A. A., Karpov, A., Frolova, O., Grigorev, A., & Lyakso, E. (2017). Emotion, age, and gender classification in children\u2019s speech by humans and machines. Computer Speech & Language, 46, 268\u2013283.","journal-title":"Computer Speech & Language"},{"key":"1929_CR41","doi-asserted-by":"publisher","first-page":"5","DOI":"10.1016\/j.knosys.2016.10.008","volume":"115","author":"Z Qawaqneh","year":"2017","unstructured":"Qawaqneh, Z., Mallouh, A. A., & Barkana, B. D. (2017). Deep neural network framework and transformed mfccs for speaker\u2019s age and gender classification. Knowledge-Based Systems, 115, 5\u201314.","journal-title":"Knowledge-Based Systems"},{"issue":"2","key":"1929_CR42","doi-asserted-by":"publisher","first-page":"101","DOI":"10.4316\/AECE.2018.02013","volume":"18","author":"O B\u00fcy\u00fck","year":"2018","unstructured":"B\u00fcy\u00fck, O., & Arslan, M. L. (2018). Combination of long-term and shortterm features for age identification from voice. Advanced Electrical Computer Engineering, 18(2), 101\u2013108.","journal-title":"Advanced Electrical Computer Engineering"},{"issue":"17","key":"1929_CR43","doi-asserted-by":"publisher","first-page":"5892","DOI":"10.3390\/s21175892","volume":"21","author":"A Tursunov","year":"2021","unstructured":"Tursunov, A., Mustaqeem, Choeh, J. Y., & Kwon, S. (2021). Age and gender recognition using a convolutional neural network with a specially designed multi-attention module through speech spectrograms. Sensors, 21(17), 5892.","journal-title":"Sensors"},{"issue":"1","key":"1929_CR44","doi-asserted-by":"publisher","first-page":"169","DOI":"10.3390\/math11010169","volume":"11","author":"D Vlaj","year":"2022","unstructured":"Vlaj, D., & Zgank, A. (2022). Acoustic gender and age classification as an aid to human-computer interaction in a smart home environment. Mathematics, 11(1), 169.","journal-title":"Mathematics"},{"issue":"3","key":"1929_CR45","doi-asserted-by":"publisher","first-page":"3535","DOI":"10.1007\/s11042-021-11614-4","volume":"81","author":"HA S\u00e1nchez-Hevia","year":"2022","unstructured":"S\u00e1nchez-Hevia, H. A., Gil-Pita, R., Utrilla-Manso, M., & Rosa-Zurera, M. (2022). Age group classification and gender recognition from speech with temporal convolutional neural networks. Multimedia Tools and Application, 81(3), 3535\u20133552.","journal-title":"Multimedia Tools and Application"},{"key":"1929_CR46","doi-asserted-by":"publisher","first-page":"3065","DOI":"10.1007\/s00521-023-09153-0","volume":"36","author":"E Y\u00fccesoy","year":"2024","unstructured":"Y\u00fccesoy, E. (2024). Speaker age and gender recognition using 1d and 2d convolutional neural networks. Neural Computing and Application, 36, 3065\u20133075.","journal-title":"Neural Computing and Application"}],"container-title":["Journal of Signal Processing Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11265-024-01929-4.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11265-024-01929-4\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11265-024-01929-4.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,1,17]],"date-time":"2025-01-17T09:40:00Z","timestamp":1737106800000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11265-024-01929-4"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,10]]},"references-count":46,"journal-issue":{"issue":"10","published-print":{"date-parts":[[2024,10]]}},"alternative-id":["1929"],"URL":"https:\/\/doi.org\/10.1007\/s11265-024-01929-4","relation":{},"ISSN":["1939-8018","1939-8115"],"issn-type":[{"value":"1939-8018","type":"print"},{"value":"1939-8115","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,10]]},"assertion":[{"value":"20 March 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"19 July 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"25 July 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"1 November 2024","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors state that they are aware of no personal or financial conflicts that might have affected the research described in this study.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no interpersonal or financial conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of Interest"}}]}}