{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,2]],"date-time":"2026-07-02T08:16:27Z","timestamp":1782980187834,"version":"3.54.5"},"reference-count":26,"publisher":"IEEE","license":[{"start":{"date-parts":[[2021,6,6]],"date-time":"2021-06-06T00:00:00Z","timestamp":1622937600000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2021,6,6]],"date-time":"2021-06-06T00:00:00Z","timestamp":1622937600000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"}],"content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2021,6,6]]},"DOI":"10.1109\/icassp39728.2021.9414292","type":"proceedings-article","created":{"date-parts":[[2021,5,13]],"date-time":"2021-05-13T19:53:45Z","timestamp":1620935625000},"page":"6254-6258","source":"Crossref","is-referenced-by-count":20,"title":["AISpeech-SJTU Accent Identification System for the Accented English Speech Recognition Challenge"],"prefix":"10.1109","author":[{"given":"Houjun","family":"Huang","sequence":"first","affiliation":[{"name":"AISpeech Ltd,Suzhou,China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xu","family":"Xiang","sequence":"additional","affiliation":[{"name":"AISpeech Ltd,Suzhou,China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yexin","family":"Yang","sequence":"additional","affiliation":[{"name":"Shanghai Jiao Tong University,MoE Key Lab of Artificial Intelligence, AI Institute SpeechLab,Department of Computer Science and Engineering,Shanghai,China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Rao","family":"Ma","sequence":"additional","affiliation":[{"name":"Shanghai Jiao Tong University,MoE Key Lab of Artificial Intelligence, AI Institute SpeechLab,Department of Computer Science and Engineering,Shanghai,China"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yanmin","family":"Qian","sequence":"additional","affiliation":[{"name":"Shanghai Jiao Tong University,MoE Key Lab of Artificial Intelligence, AI Institute SpeechLab,Department of Computer Science and Engineering,Shanghai,China"}],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"263","reference":[{"key":"ref10","article-title":"The kaldi speech recognition toolkit","author":"povey","year":"2011","journal-title":"2011 IEEE Workshop on Automatic Speech Recognition &amp; Understanding"},{"key":"ref11","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP40776.2020.9053436"},{"key":"ref12","first-page":"3171","article-title":"Fastspeech: Fast, robust and controllable text to speech","author":"ren","year":"2019","journal-title":"Advances in neural information processing systems"},{"key":"ref13","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2019.8682804"},{"key":"ref14","doi-asserted-by":"publisher","DOI":"10.21437\/Interspeech.2018-1456"},{"key":"ref15","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2018.8461368"},{"key":"ref16","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2018.8461375"},{"key":"ref17","doi-asserted-by":"publisher","DOI":"10.21437\/Interspeech.2020-2650"},{"key":"ref18","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2019.2938758"},{"key":"ref19","first-page":"7132","article-title":"Squeeze-and-excitation networks","author":"hu","year":"2018","journal-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition"},{"key":"ref4","doi-asserted-by":"publisher","DOI":"10.21437\/Interspeech.2019-2737"},{"key":"ref3","doi-asserted-by":"publisher","DOI":"10.1109\/ASRU.2001.1034657"},{"key":"ref6","article-title":"The accented english speech recognition challenge 2020: open datasets, tracks, baselines, results and methods","author":"shi","year":"2020"},{"key":"ref5","doi-asserted-by":"publisher","DOI":"10.21437\/Interspeech.2016-1148"},{"key":"ref8","doi-asserted-by":"publisher","DOI":"10.21437\/Interspeech.2019-1778"},{"key":"ref7","doi-asserted-by":"publisher","DOI":"10.21437\/Interspeech.2016-1043"},{"key":"ref2","article-title":"Improving accent identification through knowledge of english syllable structure","author":"berkling","year":"1998","journal-title":"Intl Conference on Spoken Language Processing"},{"key":"ref9","article-title":"Musan: A music, speech, and noise corpus","author":"snyder","year":"2015","journal-title":"arXiv preprint arXiv 1510 08484"},{"key":"ref1","article-title":"Foreign accent identification based on prosodic parameters","author":"piat","year":"2008","journal-title":"Ninth Annual Conference of the International Speech Communication Association"},{"key":"ref20","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00482"},{"key":"ref22","first-page":"1097","article-title":"Imagenet classification with deep convolutional neural networks","author":"krizhevsky","year":"2012","journal-title":"Advances in neural information processing systems"},{"key":"ref21","doi-asserted-by":"publisher","DOI":"10.1109\/APSIPAASC47483.2019.9023039"},{"key":"ref24","first-page":"8026","article-title":"Pytorch: An imperative style, high-performance deep learning library","author":"paszke","year":"2019","journal-title":"Advances in neural information processing systems"},{"key":"ref23","article-title":"Very deep convolutional networks for large-scale image recognition","author":"simonyan","year":"2014","journal-title":"arXiv preprint arXiv 1409 1556"},{"key":"ref26","first-page":"2579","article-title":"Visualizing high-dimensional data using t-sne","volume":"9","author":"hinton","year":"2008","journal-title":"Journal of Machine Learning Research"},{"key":"ref25","article-title":"Aispeech-sjtu asr system for the accented english speech recognition challenge","author":"tian","year":"2020"}],"event":{"name":"ICASSP 2021 - 2021 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","location":"Toronto, ON, Canada","start":{"date-parts":[[2021,6,6]]},"end":{"date-parts":[[2021,6,11]]}},"container-title":["ICASSP 2021 - 2021 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)"],"original-title":[],"link":[{"URL":"http:\/\/xplorestaging.ieee.org\/ielx7\/9413349\/9413350\/09414292.pdf?arnumber=9414292","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,12,19]],"date-time":"2024-12-19T19:18:35Z","timestamp":1734635915000},"score":1,"resource":{"primary":{"URL":"https:\/\/ieeexplore.ieee.org\/document\/9414292\/"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2021,6,6]]},"references-count":26,"URL":"https:\/\/doi.org\/10.1109\/icassp39728.2021.9414292","relation":{},"subject":[],"published":{"date-parts":[[2021,6,6]]}}}