{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2024,6,11]],"date-time":"2024-06-11T00:05:30Z","timestamp":1718064330462},"reference-count":43,"publisher":"Springer Science and Business Media LLC","issue":"8","license":[{"start":{"date-parts":[[2015,8,25]],"date-time":"2015-08-25T00:00:00Z","timestamp":1440460800000},"content-version":"tdm","delay-in-days":0,"URL":"http:\/\/www.springer.com\/tdm"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Neural Comput &amp; Applic"],"published-print":{"date-parts":[[2016,11]]},"DOI":"10.1007\/s00521-015-2030-9","type":"journal-article","created":{"date-parts":[[2015,8,24]],"date-time":"2015-08-24T06:04:44Z","timestamp":1440396284000},"page":"2615-2628","update-policy":"http:\/\/dx.doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Voice conversion system using salient sub-bands and radial basis function"],"prefix":"10.1007","volume":"27","author":[{"given":"Jagannath","family":"Nirmal","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mukesh","family":"Zaveri","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Suprava","family":"Patnaik","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Pramod","family":"Kachare","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2015,8,25]]},"reference":[{"key":"2030_CR1","doi-asserted-by":"crossref","unstructured":"Kain A, Macon MW (2001) Design and evaluation of a voice conversion algorithm based on spectral envelope mapping and residual prediction. In: Proceedings of IEEE international conference acoustics speech signal processing, vol 2, pp. 813\u2013816","DOI":"10.1109\/ICASSP.2001.941039"},{"key":"2030_CR2","doi-asserted-by":"crossref","first-page":"211","DOI":"10.1016\/S0167-6393(99)00015-1","volume":"28","author":"LM Arslan","year":"1999","unstructured":"Arslan LM (1999) Speaker transformation algorithm using segmental code books (STASC). Speech Commun 28:211\u2013226","journal-title":"Speech Commun"},{"key":"2030_CR3","doi-asserted-by":"crossref","first-page":"641","DOI":"10.1109\/TASL.2006.876760","volume":"15","author":"K Lee","year":"2007","unstructured":"Lee K (2007) Statistical approach for voice personality transformation. IEEE Trans Audio Speech Lang Process 15:641\u2013651","journal-title":"IEEE Trans Audio Speech Lang Process"},{"key":"2030_CR4","volume-title":"Fundamentals of speech recognition","author":"L Rabiner","year":"1993","unstructured":"Rabiner L, Juang BH (1993) Fundamentals of speech recognition. Prentice Hall of India, New Delhi"},{"key":"2030_CR5","doi-asserted-by":"crossref","unstructured":"Furui S (1986) Research on individuality features in speech waves and automatic speaker recognition techniques. Speech Commun 5(2):183\u2013197","DOI":"10.1016\/0167-6393(86)90007-5"},{"issue":"3","key":"2030_CR6","doi-asserted-by":"crossref","first-page":"474","DOI":"10.1016\/j.csl.2009.03.003","volume":"24","author":"KS Rao","year":"2010","unstructured":"Rao KS (2010) Voice conversion by mapping the speaker-specific features using pitch synchronous approach. Comput Speech Lang Process 24(3):474\u2013494","journal-title":"Comput Speech Lang Process"},{"key":"2030_CR7","doi-asserted-by":"crossref","first-page":"36","DOI":"10.1155\/S1110865701000117","volume":"1","author":"C Drioli","year":"2001","unstructured":"Drioli C (2001) Radial basis function networks for conversion of sound speech spectra. EURASIP J Appl Signal Process 1:36\u201340","journal-title":"EURASIP J Appl Signal Process"},{"key":"2030_CR8","doi-asserted-by":"crossref","unstructured":"Strang G, Nguyen T (1997) Wavelets and filter banks. Wellesley Cambridge Press, Wellesley","DOI":"10.1093\/oso\/9780195094237.003.0002"},{"key":"2030_CR9","first-page":"145","volume":"1","author":"H Valbret","year":"1992","unstructured":"Valbret H, Moulines E, Tubach JP (1992) Voice transformation using PSOLA technique. Speech Commun 1:145\u2013148","journal-title":"Speech Commun"},{"issue":"1","key":"2030_CR10","first-page":"17","volume":"2","author":"AN Chadha","year":"2014","unstructured":"Chadha AN, Nirmal JH, Kachare P (2014) A Comparative performance of various speech analysis-synthesis techniques. Int J Signal Process Syst 2(1):17","journal-title":"Int J Signal Process Syst"},{"key":"2030_CR11","unstructured":"Deshpande Mangesh S, Holambe Raghunath S (2010) Speaker identification using admissible wavelet packet based decomposition. World Acad Sci Eng Technol 37:736\u2013739"},{"key":"2030_CR12","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1016\/j.asoc.2014.06.040","volume":"24","author":"JH Nirmal","year":"2014","unstructured":"Nirmal JH, Zaveri M, Patnaik S, Kachare P (2014) Voice conversion using general egression neural network. Appl Soft Comput 24:1\u201312","journal-title":"Appl Soft Comput"},{"key":"2030_CR13","doi-asserted-by":"crossref","unstructured":"Narendranath M, Murthy HA, Rajendran S, Yegnanarayana B (1995) Transformation of formants for voice conversion using artificial neural networks. Speech Commun 16(2):207\u2013216","DOI":"10.1016\/0167-6393(94)00058-I"},{"key":"2030_CR14","unstructured":"Stylianou Y (1996) Harmonic plus noise models for speech, combined with statistical methods, for speech and speaker modification, Ph.D. dissertation, cole Nationale Superieure Des Tlcommunications. Paris, France"},{"key":"2030_CR15","doi-asserted-by":"crossref","unstructured":"Stylianou Y, Capp O, Moulines E (1998) Continuous probabilistic transform for voice conversion. In: Proceedings IEEE international conference acoustics, speech, signal process. vol 6, pp. 131\u2013142","DOI":"10.1109\/89.661472"},{"key":"2030_CR16","unstructured":"Nirmal JH, Zaveri M, Patnaik S, Kachare P (2014) Complex cepstrum based voice conversion using radial basis function neural network. In: ISRN signal processing, vol 2014. Hindawi Publishing Corporation, Article ID 357048"},{"key":"2030_CR17","unstructured":"Kominek J, Black AW (2004) The CMU ARCTIC speech databases. In: Proceedings 5th ISCA speech synthesis workshop (SSW5), Pittsburgh, PA, pp. 223\u2013224"},{"key":"2030_CR18","doi-asserted-by":"crossref","unstructured":"Guidoa Rodrigo C, Vieiraa Lucimar Sasso, Juniora Sylvio Barbon (2007) A neural wavelet architectures for voice conversion. Sci Direct Neurocomput 71:174\u2013180","DOI":"10.1016\/j.neucom.2007.08.010"},{"key":"2030_CR19","doi-asserted-by":"crossref","first-page":"165","DOI":"10.1016\/0167-6393(94)00053-D","volume":"16","author":"H Kuwabura","year":"1995","unstructured":"Kuwabura H, Sagisaka Y (1995) Acoustic characteristics of speaker individuality: control and conversion. Speech Commun 16:165\u2013173","journal-title":"Speech Commun"},{"key":"2030_CR20","doi-asserted-by":"crossref","first-page":"3332","DOI":"10.1016\/j.asoc.2012.05.027","volume":"12","author":"RH Laskara","year":"2012","unstructured":"Laskara RH, Chakrabartyb D, Talukdara FA, Sreenivasa Raoc K, Banerjeea K (2012) Comparing ANN and GMM in a voice conversion framework. Appl Soft Comput 12:3332\u20133342","journal-title":"Appl Soft Comput"},{"key":"2030_CR21","doi-asserted-by":"crossref","first-page":"131142","DOI":"10.1109\/89.661472","volume":"6","author":"Y Stylianou","year":"1998","unstructured":"Stylianou Y, Cappe Y, Moulines E (1998) Continuous probabilistic transform for voice conversion. IEEE Trans Speech Audio Process 6:131142","journal-title":"IEEE Trans Speech Audio Process"},{"issue":"5","key":"2030_CR22","doi-asserted-by":"crossref","first-page":"954","DOI":"10.1109\/TASL.2010.2047683","volume":"18","author":"S Desai","year":"2010","unstructured":"Desai S, Black AW, Yegnanarayana B, Prahallad K (2010) Spectral mapping using artificial neural networks for voice conversion. IEEE Trans Audio Speech Lang Process 18(5):954\u2013964","journal-title":"IEEE Trans Audio Speech Lang Process"},{"key":"2030_CR23","doi-asserted-by":"crossref","unstructured":"Childers DG, Yegnanarayana B, Wu K (1985) Voice conversion: factor responsible for quality. In: Proceedings of IEEE ICASSP, pp. 530\u2013533","DOI":"10.1109\/ICASSP.1985.1168479"},{"key":"2030_CR24","doi-asserted-by":"crossref","unstructured":"Kuwabura H, Sagisaka Y (1995) Acoustic characteristics of speaker individuality: control and conversion. Speech Commun 16:165\u2013173","DOI":"10.1016\/0167-6393(94)00053-D"},{"issue":"2","key":"2030_CR25","doi-asserted-by":"crossref","first-page":"207","DOI":"10.1016\/0167-6393(94)00058-I","volume":"16","author":"M Narendranath","year":"1995","unstructured":"Narendranath M, Murthy HA, Rajendran S, Yegnanarayana B (1995) Transformation of formants for voice conversion using artificial neural networks. Speech Commun 16(2):207\u2013216","journal-title":"Speech Commun"},{"key":"2030_CR26","doi-asserted-by":"crossref","unstructured":"Nirmal JH, Zaveri M, Patnaik S, Kachare P (2013) A novel voice conversion approach using admissible wavelet packet decomposition. EURASIP J Audio Speech Music Process 2013:28","DOI":"10.1186\/1687-4722-2013-28"},{"issue":"5","key":"2030_CR27","doi-asserted-by":"crossref","first-page":"912","DOI":"10.1109\/TASL.2010.2041699","volume":"18","author":"E Helander","year":"2010","unstructured":"Helander E, Virtanen T, Jani N, Gabbouj M (2010) Voice conversion using partial least squares regression. IEEE Trans Audio Speech Lang Process 18(5):912\u2013921","journal-title":"IEEE Trans Audio Speech Lang Process"},{"key":"2030_CR28","doi-asserted-by":"crossref","unstructured":"Desai S, Black AW, Yegnanarayana B, Prahallad K (2010) Spectral mapping using artificial neural networks for voice conversion. IEEE Trans Audio Speech Lang Process 18(5):954\u2013964","DOI":"10.1109\/TASL.2010.2047683"},{"key":"2030_CR29","doi-asserted-by":"crossref","unstructured":"Masuko T, Tokuda K, Kobayashi T, Imai S (1996) Speech synthesis using HMMS with dynamic features. In: Proceedings IEEE international conference acoustics, speech, signal processing, pp. 389\u2013392","DOI":"10.1109\/ICASSP.1996.541114"},{"issue":"3","key":"2030_CR30","first-page":"3297","volume":"10","author":"C Orphanidou","year":"2004","unstructured":"Orphanidou C, Moroz IM, Roberts SJ (2004) Wavelet-based voice morphing. WSEAS J Syst 10(3):3297\u20133302","journal-title":"WSEAS J Syst"},{"key":"2030_CR31","doi-asserted-by":"crossref","first-page":"174","DOI":"10.1016\/j.neucom.2007.08.010","volume":"71","author":"Rodrigo C Guidoa","year":"2007","unstructured":"Guidoa Rodrigo C, Vieiraa Lucimar Sasso, Juniora Sylvio Barbon (2007) A neural wavelet architectures for voice conversion. Sci Direct Neurocomput 71:174\u2013180","journal-title":"Sci Direct Neurocomput"},{"key":"2030_CR32","unstructured":"Nirmal JH, Patnaik SS, Zaveri MA (2012) Voice transformation using radial basis function. In: Third international conference on recent trends in information. Telecommunication and computing ITC 2012, Springer, Berlin, pp. 271\u2013276"},{"issue":"2","key":"2030_CR33","doi-asserted-by":"crossref","first-page":"183","DOI":"10.1016\/0167-6393(86)90007-5","volume":"5","author":"S Furui","year":"1986","unstructured":"Furui S (1986) Research on individuality features in speech waves and automatic speaker recognition techniques. Speech Commun 5(2):183\u2013197","journal-title":"Speech Commun"},{"key":"2030_CR34","volume-title":"Wavelets and filter banks","author":"G Strang","year":"1997","unstructured":"Strang G, Nguyen T (1997) Wavelets and filter banks. Wellesley Cambridge Press, Wellesley"},{"key":"2030_CR35","doi-asserted-by":"crossref","unstructured":"Rao KS (2010) Voice conversion by mapping the speaker-specific features using pitch synchronous approach. Comput Speech Lang Process 24(3):474\u2013494","DOI":"10.1016\/j.csl.2009.03.003"},{"key":"2030_CR36","first-page":"736","volume":"37","author":"Mangesh S Deshpande","year":"2010","unstructured":"Deshpande Mangesh S, Holambe Raghunath S (2010) Speaker identification using admissible wavelet packet based decomposition. World Acad Sci Eng Technol 37:736\u2013739","journal-title":"World Acad Sci Eng Technol"},{"key":"2030_CR37","doi-asserted-by":"crossref","first-page":"28","DOI":"10.1186\/1687-4722-2013-28","volume":"2013","author":"JH Nirmal","year":"2013","unstructured":"Nirmal JH, Zaveri M, Patnaik S, Kachare P (2013) A novel voice conversion approach using admissible wavelet packet decomposition. EURASIP J Audio Speech Music Process 2013:28","journal-title":"EURASIP J Audio Speech Music Process"},{"key":"2030_CR38","doi-asserted-by":"crossref","unstructured":"Xugang Lu, Dang Jianwu (2008) An investigation of dependencies between frequency components and speaker characteristics for text-independent speaker identification. Speech Commun 50:312\u2013322","DOI":"10.1016\/j.specom.2007.10.005"},{"issue":"3","key":"2030_CR39","doi-asserted-by":"crossref","first-page":"379","DOI":"10.1002\/j.1538-7305.1948.tb01338.x","volume":"27","author":"CE Shannon","year":"1948","unstructured":"Shannon CE (1948) A mathematical theory of communication. Bell Syst Tech J 27(3):379\u2013423","journal-title":"Bell Syst Tech J"},{"issue":"1","key":"2030_CR40","doi-asserted-by":"crossref","first-page":"3","DOI":"10.1145\/584091.584093","volume":"5","author":"CE Shannon","year":"2001","unstructured":"Shannon CE (2001) A mathematical theory of communication. ACM SIGMOBILE Mob Comput Commun Rev 5(1):3\u201355","journal-title":"ACM SIGMOBILE Mob Comput Commun Rev"},{"issue":"6","key":"2030_CR41","doi-asserted-by":"crossref","first-page":"066138","DOI":"10.1103\/PhysRevE.69.066138","volume":"69","author":"A Kraskov","year":"2004","unstructured":"Kraskov A, St\u00f6gbauer H, Grassberger P (2004) Estimating mutual information. Phys Rev E 69(6):066138","journal-title":"Phys Rev E"},{"key":"2030_CR42","doi-asserted-by":"crossref","first-page":"312","DOI":"10.1016\/j.specom.2007.10.005","volume":"50","author":"Lu Xugang","year":"2008","unstructured":"Xugang Lu, Dang Jianwu (2008) An investigation of dependencies between frequency components and speaker characteristics for text-independent speaker identification. Speech Commun 50:312\u2013322","journal-title":"Speech Commun"},{"key":"2030_CR43","unstructured":"Reza Fazlollah M (1961, 1994) An introduction to information theory. Dover Publications Inc., New York"}],"container-title":["Neural Computing and Applications"],"original-title":[],"language":"en","link":[{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s00521-015-2030-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"http:\/\/link.springer.com\/article\/10.1007\/s00521-015-2030-9\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s00521-015-2030-9.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"},{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s00521-015-2030-9","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,6,10]],"date-time":"2024-06-10T18:06:58Z","timestamp":1718042818000},"score":1,"resource":{"primary":{"URL":"http:\/\/link.springer.com\/10.1007\/s00521-015-2030-9"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2015,8,25]]},"references-count":43,"journal-issue":{"issue":"8","published-print":{"date-parts":[[2016,11]]}},"alternative-id":["2030"],"URL":"https:\/\/doi.org\/10.1007\/s00521-015-2030-9","relation":{},"ISSN":["0941-0643","1433-3058"],"issn-type":[{"value":"0941-0643","type":"print"},{"value":"1433-3058","type":"electronic"}],"subject":[],"published":{"date-parts":[[2015,8,25]]}}}