{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,2]],"date-time":"2026-07-02T16:41:00Z","timestamp":1783010460899,"version":"3.54.6"},"publisher-location":"New York, NY, USA","reference-count":80,"publisher":"ACM","license":[{"start":{"date-parts":[[2022,11,7]],"date-time":"2022-11-07T00:00:00Z","timestamp":1667779200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.acm.org\/publications\/policies\/copyright_policy#Background"}],"funder":[{"DOI":"10.13039\/100000001","name":"NSF (National Science Foundation)","doi-asserted-by":"publisher","award":["IIS-2046590 CNS-2114411 CNS-1942610 CNS-2114407"],"award-info":[{"award-number":["IIS-2046590 CNS-2114411 CNS-1942610 CNS-2114407"]}],"id":[{"id":"10.13039\/100000001","id-type":"DOI","asserted-by":"publisher"}]},{"name":"UK's National Research Centre on Privacy Harm Reduction and Adversarial Influence Online","award":["EP\/V011189\/1"],"award-info":[{"award-number":["EP\/V011189\/1"]}]},{"DOI":"10.13039\/501100009318","name":"Helmholtz Association","doi-asserted-by":"publisher","award":["ZT-I-OO1 4"],"award-info":[{"award-number":["ZT-I-OO1 4"]}],"id":[{"id":"10.13039\/501100009318","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2022,11,7]]},"DOI":"10.1145\/3548606.3560599","type":"proceedings-article","created":{"date-parts":[[2022,11,7]],"date-time":"2022-11-07T11:41:28Z","timestamp":1667821288000},"page":"2659-2673","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":39,"title":["Why So Toxic?"],"prefix":"10.1145","author":[{"given":"Wai Man","family":"Si","sequence":"first","affiliation":[{"name":"CISPA Helmholtz Center for Information Security, Saarbruceken, Germany"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Michael","family":"Backes","sequence":"additional","affiliation":[{"name":"CISPA Helmholtz Center for Information Security, Saarbruceken, Germany"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jeremy","family":"Blackburn","sequence":"additional","affiliation":[{"name":"Binghamton University, Binghamton, NY, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Emiliano","family":"De Cristofaro","sequence":"additional","affiliation":[{"name":"University College London, London, United Kingdom"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Gianluca","family":"Stringhini","sequence":"additional","affiliation":[{"name":"Boston University, Boston, MA, USA"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Savvas","family":"Zannettou","sequence":"additional","affiliation":[{"name":"TU Delft, Delft, Netherlands"}],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yang","family":"Zhang","sequence":"additional","affiliation":[{"name":"CISPA Helmholtz Center for Information Security, Saarbruceken, Germany"}],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"320","published-online":{"date-parts":[[2022,11,7]]},"reference":[{"key":"e_1_3_2_2_1_1","unstructured":"https:\/\/www.perspectiveapi.com.  https:\/\/www.perspectiveapi.com."},{"key":"e_1_3_2_2_2_1","unstructured":"https:\/\/github.com\/unitaryai\/detoxify.  https:\/\/github.com\/unitaryai\/detoxify."},{"key":"e_1_3_2_2_3_1","unstructured":"https:\/\/huggingface.co\/ykilcher\/gpt-4chan.  https:\/\/huggingface.co\/ykilcher\/gpt-4chan."},{"key":"e_1_3_2_2_4_1","unstructured":"https:\/\/woebothealth.com\/.  https:\/\/woebothealth.com\/."},{"key":"e_1_3_2_2_5_1","doi-asserted-by":"publisher","DOI":"10.1207\/s15516709cog0901_7"},{"key":"e_1_3_2_2_6_1","first-page":"830","volume-title":"Jeremy Blackburn. The Pushshift Reddit Dataset. In International Conference on Web and Social Media (ICWSM)","author":"Baumgartner Jason","year":"2020","unstructured":"Jason Baumgartner , Savvas Zannettou , Brian Keegan , Megan Squire , and Jeremy Blackburn. The Pushshift Reddit Dataset. In International Conference on Web and Social Media (ICWSM) , pages 830 -- 839 . AAAI, 2020 . Jason Baumgartner, Savvas Zannettou, Brian Keegan, Megan Squire, and Jeremy Blackburn. The Pushshift Reddit Dataset. In International Conference on Web and Social Media (ICWSM), pages 830--839. AAAI, 2020."},{"key":"e_1_3_2_2_7_1","first-page":"932","volume-title":"Pascal Vincent. A Neural Probabilistic Language Model. In Annual Conference on Neural Information Processing Systems (NIPS)","author":"Bengio Yoshua","year":"2000","unstructured":"Yoshua Bengio , R\u00e9jean Ducharme , and Pascal Vincent. A Neural Probabilistic Language Model. In Annual Conference on Neural Information Processing Systems (NIPS) , pages 932 -- 938 . NIPS, 2000 . Yoshua Bengio, R\u00e9jean Ducharme, and Pascal Vincent. A Neural Probabilistic Language Model. In Annual Conference on Neural Information Processing Systems (NIPS), pages 932--938. NIPS, 2000."},{"key":"e_1_3_2_2_8_1","volume-title":"Jason Weston. Learning End-to-End Goal- Oriented Dialog. In International Conference on Learning Representations (ICLR)","author":"Bordes Antoine","year":"2017","unstructured":"Antoine Bordes , Y- Lan Boureau , and Jason Weston. Learning End-to-End Goal- Oriented Dialog. In International Conference on Learning Representations (ICLR) , 2017 . Antoine Bordes, Y-Lan Boureau, and Jason Weston. Learning End-to-End Goal- Oriented Dialog. In International Conference on Learning Representations (ICLR), 2017."},{"key":"e_1_3_2_2_9_1","first-page":"148","volume-title":"ACM Conference on Web Science (WebSci)","author":"Chandra Mohit","year":"2021","unstructured":"Mohit Chandra , Dheeraj Reddy Pailla , Himanshu Bhatia , AadilMehdi J. Sanchawala , Manish Gupta , Manish Shrivastava , and Ponnurangam Kumaraguru . ? Subverting the Jewtocracy\" : Online Antisemitism Detection Using Multimodal Deep Learning . In ACM Conference on Web Science (WebSci) , pages 148 -- 157 . ACM, 2021 . Mohit Chandra, Dheeraj Reddy Pailla, Himanshu Bhatia, AadilMehdi J. Sanchawala, Manish Gupta, Manish Shrivastava, and Ponnurangam Kumaraguru. ?Subverting the Jewtocracy\": Online Antisemitism Detection Using Multimodal Deep Learning. In ACM Conference on Web Science (WebSci), pages 148--157. ACM, 2021."},{"key":"e_1_3_2_2_10_1","doi-asserted-by":"publisher","DOI":"10.1145\/3465336.3475111"},{"key":"e_1_3_2_2_11_1","volume-title":"Gianluca Stringhini, Athena Vakali, and Nicolas Kourtellis. Detecting Cyberbullying and Cyberaggression in Social Media. CoRR abs\/1907.08873","author":"Chatzakou Despoina","year":"2019","unstructured":"Despoina Chatzakou , Ilias Leontiadis , Jeremy Blackburn , Emiliano De Cristofaro , Gianluca Stringhini, Athena Vakali, and Nicolas Kourtellis. Detecting Cyberbullying and Cyberaggression in Social Media. CoRR abs\/1907.08873 , 2019 . Despoina Chatzakou, Ilias Leontiadis, Jeremy Blackburn, Emiliano De Cristofaro, Gianluca Stringhini, Athena Vakali, and Nicolas Kourtellis. Detecting Cyberbullying and Cyberaggression in Social Media. CoRR abs\/1907.08873, 2019."},{"key":"e_1_3_2_2_12_1","doi-asserted-by":"publisher","DOI":"10.1145\/3166054.3166058"},{"key":"e_1_3_2_2_13_1","first-page":"512","volume-title":"Ingmar Weber. Automated Hate Speech Detection and the Problem of Offensive Language. In International Conference on Web and Social Media (ICWSM)","author":"Davidson Thomas","year":"2017","unstructured":"Thomas Davidson , Dana Warmsley , Michael W. Macy , and Ingmar Weber. Automated Hate Speech Detection and the Problem of Offensive Language. In International Conference on Web and Social Media (ICWSM) , pages 512 -- 515 . AAAI, 2017 . Thomas Davidson, Dana Warmsley, Michael W. Macy, and Ingmar Weber. Automated Hate Speech Detection and the Problem of Offensive Language. In International Conference on Web and Social Media (ICWSM), pages 512--515. AAAI, 2017."},{"key":"e_1_3_2_2_14_1","first-page":"4171","volume-title":"Conference of the North American","author":"Devlin Jacob","year":"2019","unstructured":"Jacob Devlin , Ming-Wei Chang , Kenton Lee , and Kristina Toutanova . BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding . In Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (NAACL-HLT), pages 4171 -- 4186 . ACL , 2019 . Jacob Devlin, Ming-Wei Chang, Kenton Lee, and Kristina Toutanova. BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding. In Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (NAACL-HLT), pages 4171--4186. ACL, 2019."},{"key":"e_1_3_2_2_15_1","volume-title":"Anticipating Safety Issues in E2E Conversational AI: Framework and Tooling. CoRR abs\/2107.03451","author":"Dinan Emily","year":"2021","unstructured":"Emily Dinan , Gavin Abercrombie , A. Stevie Bergman , Shannon L. Spruit , Dirk Hovy , Y- Lan Boureau , and Verena Rieser . Anticipating Safety Issues in E2E Conversational AI: Framework and Tooling. CoRR abs\/2107.03451 , 2021 . Emily Dinan, Gavin Abercrombie, A. Stevie Bergman, Shannon L. Spruit, Dirk Hovy, Y-Lan Boureau, and Verena Rieser. Anticipating Safety Issues in E2E Conversational AI: Framework and Tooling. CoRR abs\/2107.03451, 2021."},{"key":"e_1_3_2_2_16_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2020.emnlp-main.656"},{"key":"e_1_3_2_2_17_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/D19-1461"},{"key":"e_1_3_2_2_18_1","volume-title":"Beam Search Strategies for Neural Machine Translation. CoRR abs\/1702.01806","author":"Freitag Markus","year":"2017","unstructured":"Markus Freitag and Yaser Al-Onaizan . Beam Search Strategies for Neural Machine Translation. CoRR abs\/1702.01806 , 2017 . Markus Freitag and Yaser Al-Onaizan. Beam Search Strategies for Neural Machine Translation. CoRR abs\/1702.01806, 2017."},{"key":"e_1_3_2_2_19_1","first-page":"2","volume-title":"Annual Meeting of the Association for Computational Linguistics (ACL)","author":"Gao Jianfeng","year":"2018","unstructured":"Jianfeng Gao , Michel Galley , and Lihong Li . Neural Approaches to Conversational AI . In Annual Meeting of the Association for Computational Linguistics (ACL) , pages 2 -- 7 . ACL, 2018 . Jianfeng Gao, Michel Galley, and Lihong Li. Neural Approaches to Conversational AI. In Annual Meeting of the Association for Computational Linguistics (ACL), pages 2--7. ACL, 2018."},{"key":"e_1_3_2_2_20_1","volume-title":"RealToxicityPrompts: Evaluating Neural Toxic Degeneration in Language Models. CoRR abs\/2009.11462","author":"Gehman Samuel","year":"2020","unstructured":"Samuel Gehman , Suchin Gururangan , Maarten Sap , Yejin Choi , and Noah A. Smith . RealToxicityPrompts: Evaluating Neural Toxic Degeneration in Language Models. CoRR abs\/2009.11462 , 2020 . Samuel Gehman, Suchin Gururangan, Maarten Sap, Yejin Choi, and Noah A. Smith. RealToxicityPrompts: Evaluating Neural Toxic Degeneration in Language Models. CoRR abs\/2009.11462, 2020."},{"key":"e_1_3_2_2_21_1","volume-title":"He and James R. Glass. Detecting Egregious Responses in Neural Sequence-to-sequence Models. In International Conference on Learning Representations (ICLR)","author":"Tianxing","year":"2019","unstructured":"Tianxing He and James R. Glass. Detecting Egregious Responses in Neural Sequence-to-sequence Models. In International Conference on Learning Representations (ICLR) , 2019 . Tianxing He and James R. Glass. Detecting Egregious Responses in Neural Sequence-to-sequence Models. In International Conference on Learning Representations (ICLR), 2019."},{"key":"e_1_3_2_2_22_1","first-page":"92","volume-title":"International Conference on Web and Social Media (ICWSM)","author":"Hine Gabriel Emile","year":"2017","unstructured":"Gabriel Emile Hine , Jeremiah Onaolapo , Emiliano De Cristofaro , Nicolas Kourtel- lis, Ilias Leontiadis , Riginos Samaras , Gianluca Stringhini , and Jeremy Blackburn . Kek, Cucks, and God Emperor Trump : A Measurement Study of 4chan's Politi- cally Incorrect Forum and Its Effects on the Web . In International Conference on Web and Social Media (ICWSM) , pages 92 -- 101 . AAAI, 2017 . Gabriel Emile Hine, Jeremiah Onaolapo, Emiliano De Cristofaro, Nicolas Kourtel- lis, Ilias Leontiadis, Riginos Samaras, Gianluca Stringhini, and Jeremy Blackburn. Kek, Cucks, and God Emperor Trump: A Measurement Study of 4chan's Politi- cally Incorrect Forum and Its Effects on the Web. In International Conference on Web and Social Media (ICWSM), pages 92--101. AAAI, 2017."},{"key":"e_1_3_2_2_23_1","volume-title":"Distilling the Knowledge in a Neural Network. CoRR abs\/1503.02531","author":"Hinton Geoffrey E.","year":"2015","unstructured":"Geoffrey E. Hinton , Oriol Vinyals , and Jeffrey Dean . Distilling the Knowledge in a Neural Network. CoRR abs\/1503.02531 , 2015 . Geoffrey E. Hinton, Oriol Vinyals, and Jeffrey Dean. Distilling the Knowledge in a Neural Network. CoRR abs\/1503.02531, 2015."},{"key":"e_1_3_2_2_24_1","volume-title":"Yejin Choi. The Curious Case of Neural Text Degeneration. In International Conference on Learning Representations (ICLR)","author":"Holtzman Ari","year":"2020","unstructured":"Ari Holtzman , Jan Buys , Li Du , Maxwell Forbes , and Yejin Choi. The Curious Case of Neural Text Degeneration. In International Conference on Learning Representations (ICLR) , 2020 . Ari Holtzman, Jan Buys, Li Du, Maxwell Forbes, and Yejin Choi. The Curious Case of Neural Text Degeneration. In International Conference on Learning Representations (ICLR), 2020."},{"key":"e_1_3_2_2_25_1","volume-title":"Deceiving Google's Perspective API Built for Detecting Toxic Comments. CoRR abs\/1702.08138","author":"Hosseini Hossein","year":"2017","unstructured":"Hossein Hosseini , Sreeram Kannan , Baosen Zhang , and Radha Poovendran . Deceiving Google's Perspective API Built for Detecting Toxic Comments. CoRR abs\/1702.08138 , 2017 . Hossein Hosseini, Sreeram Kannan, Baosen Zhang, and Radha Poovendran. Deceiving Google's Perspective API Built for Detecting Toxic Comments. CoRR abs\/1702.08138, 2017."},{"key":"e_1_3_2_2_26_1","volume-title":"Advancing the State of the Art in Open Domain Dialog Systems through the Alexa Prize. CoRR abs\/1812.10757","author":"Khatri Chandra","year":"2018","unstructured":"Chandra Khatri , Behnam Hedayatnia , Anu Venkatesh , Jeff Nunn , Yi Pan , Qing Liu , Han Song , Anna Gottardi , Sanjeev Kwatra , Sanju Pancholi , Ming Cheng , Qinglang Chen , Lauren Stubel , Karthik Gopalakrishnan , Kate Bland , Raefer Gabriel , Arindam Mandal , Dilek Hakkani-T\u00fcr , Gene Hwang , Nate Michel , Eric King , and Rohit Prasad . Advancing the State of the Art in Open Domain Dialog Systems through the Alexa Prize. CoRR abs\/1812.10757 , 2018 . Chandra Khatri, Behnam Hedayatnia, Anu Venkatesh, Jeff Nunn, Yi Pan, Qing Liu, Han Song, Anna Gottardi, Sanjeev Kwatra, Sanju Pancholi, Ming Cheng, Qinglang Chen, Lauren Stubel, Karthik Gopalakrishnan, Kate Bland, Raefer Gabriel, Arindam Mandal, Dilek Hakkani-T\u00fcr, Gene Hwang, Nate Michel, Eric King, and Rohit Prasad. Advancing the State of the Art in Open Domain Dialog Systems through the Alexa Prize. CoRR abs\/1812.10757, 2018."},{"key":"e_1_3_2_2_27_1","volume-title":"Automatic Text Summarization of COVID-19 Medical Research Articles using BERT and GPT-2. CoRR abs\/2006.01997","author":"Kieuvongngam Virapat","year":"2020","unstructured":"Virapat Kieuvongngam , Bowen Tan , and Yiming Niu . Automatic Text Summarization of COVID-19 Medical Research Articles using BERT and GPT-2. CoRR abs\/2006.01997 , 2020 . Virapat Kieuvongngam, Bowen Tan, and Yiming Niu. Automatic Text Summarization of COVID-19 Medical Research Articles using BERT and GPT-2. CoRR abs\/2006.01997, 2020."},{"key":"e_1_3_2_2_28_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/N16-1014"},{"key":"e_1_3_2_2_29_1","volume-title":"Fast Diverse Decoding Algorithm for Neural Generation. CoRR abs\/1611.08562","author":"Li Jiwei","year":"2016","unstructured":"Jiwei Li , Will Monroe , and Dan Jurafsky . A Simple , Fast Diverse Decoding Algorithm for Neural Generation. CoRR abs\/1611.08562 , 2016 . Jiwei Li, Will Monroe, and Dan Jurafsky. A Simple, Fast Diverse Decoding Algorithm for Neural Generation. CoRR abs\/1611.08562, 2016."},{"key":"e_1_3_2_2_30_1","first-page":"1192","volume-title":"Jianfeng Gao. Deep Reinforcement Learning for Dialogue Generation. In Conference on Empirical Methods in Natural Language Processing (EMNLP)","author":"Li Jiwei","year":"2016","unstructured":"Jiwei Li , Will Monroe , Alan Ritter , Dan Jurafsky , Michel Galley , and Jianfeng Gao. Deep Reinforcement Learning for Dialogue Generation. In Conference on Empirical Methods in Natural Language Processing (EMNLP) , pages 1192 -- 1202 . ACL, 2016 . Jiwei Li, Will Monroe, Alan Ritter, Dan Jurafsky, Michel Galley, and Jianfeng Gao. Deep Reinforcement Learning for Dialogue Generation. In Conference on Empirical Methods in Natural Language Processing (EMNLP), pages 1192--1202. ACL, 2016."},{"key":"e_1_3_2_2_31_1","volume-title":"Say What I Want: Towards the Dark Side of Neural Dialogue Models. CoRR abs\/1909.06044","author":"Liu Haochen","year":"2019","unstructured":"Haochen Liu , Tyler Derr , Zitao Liu , and Jiliang Tang . Say What I Want: Towards the Dark Side of Neural Dialogue Models. CoRR abs\/1909.06044 , 2019 . Haochen Liu, Tyler Derr, Zitao Liu, and Jiliang Tang. Say What I Want: Towards the Dark Side of Neural Dialogue Models. CoRR abs\/1909.06044, 2019."},{"key":"e_1_3_2_2_32_1","volume-title":"South Korean AI chatbot pulled from Facebook after hate speech towards minorities. https:\/\/www.theguardian.com\/world\/2021\/jan\/14\/time-to-properly-socialise-hate-speech-ai-chatbot-pulled-from-facebook","author":"McCurry Justin","year":"2021","unstructured":"Justin McCurry . South Korean AI chatbot pulled from Facebook after hate speech towards minorities. https:\/\/www.theguardian.com\/world\/2021\/jan\/14\/time-to-properly-socialise-hate-speech-ai-chatbot-pulled-from-facebook , 2021 . Justin McCurry. South Korean AI chatbot pulled from Facebook after hate speech towards minorities. https:\/\/www.theguardian.com\/world\/2021\/jan\/14\/time-to-properly-socialise-hate-speech-ai-chatbot-pulled-from-facebook, 2021."},{"key":"e_1_3_2_2_33_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2022.naacl-main.204"},{"key":"e_1_3_2_2_34_1","doi-asserted-by":"publisher","DOI":"10.1145\/3457607"},{"key":"e_1_3_2_2_35_1","volume-title":"ParlAI: A Dialog Research Software Platform. CoRR abs\/1705.06476","author":"Miller Alexander H.","year":"2017","unstructured":"Alexander H. Miller , Will Feng , Adam Fisch , Jiasen Lu , Dhruv Batra , Antoine Bordes , Devi Parikh , and Jason Weston . ParlAI: A Dialog Research Software Platform. CoRR abs\/1705.06476 , 2017 . Alexander H. Miller, Will Feng, Adam Fisch, Jiasen Lu, Dhruv Batra, Antoine Bordes, Devi Parikh, and Jason Weston. ParlAI: A Dialog Research Software Platform. CoRR abs\/1705.06476, 2017."},{"key":"e_1_3_2_2_36_1","first-page":"452","volume-title":"International Conference on Web and Social Media (ICWSM)","author":"Mittos Alexandros","year":"2020","unstructured":"Alexandros Mittos , Savvas Zannettou , Jeremy Blackburn , and Emiliano De Cristo- faro. \" And We Will Fight For Our Race!\" A Measurement Study of Genetic Testing Conversations on Reddit and 4chan. In International Conference on Web and Social Media (ICWSM) , pages 452 -- 463 . AAAI, 2020 . Alexandros Mittos, Savvas Zannettou, Jeremy Blackburn, and Emiliano De Cristo- faro. \"And We Will Fight For Our Race!\" A Measurement Study of Genetic Testing Conversations on Reddit and 4chan. In International Conference on Web and Social Media (ICWSM), pages 452--463. AAAI, 2020."},{"key":"e_1_3_2_2_37_1","volume-title":"Microsoft's fun millennial AI bot, into a genocidal maniac. https:\/\/wapo.st\/3mRIOww","author":"Ohlheiser Abby","year":"2016","unstructured":"Abby Ohlheiser . Trolls turned Tay , Microsoft's fun millennial AI bot, into a genocidal maniac. https:\/\/wapo.st\/3mRIOww , 2016 . Abby Ohlheiser. Trolls turned Tay, Microsoft's fun millennial AI bot, into a genocidal maniac. https:\/\/wapo.st\/3mRIOww, 2016."},{"key":"e_1_3_2_2_38_1","first-page":"4262","volume-title":"Dit-Yan Yeung. Probing Toxic Content in Large Pre-Trained Language Models. In Annual Meeting of the Association for Computational Linguistics (ACL)","author":"Ousidhoum Nedjma","year":"2021","unstructured":"Nedjma Ousidhoum , Xinran Zhao , Tianqing Fang , Yangqiu Song , and Dit-Yan Yeung. Probing Toxic Content in Large Pre-Trained Language Models. In Annual Meeting of the Association for Computational Linguistics (ACL) , pages 4262 -- 4274 . ACL, 2021 . Nedjma Ousidhoum, Xinran Zhao, Tianqing Fang, Yangqiu Song, and Dit-Yan Yeung. Probing Toxic Content in Large Pre-Trained Language Models. In Annual Meeting of the Association for Computational Linguistics (ACL), pages 4262--4274. ACL, 2021."},{"key":"e_1_3_2_2_39_1","doi-asserted-by":"publisher","DOI":"10.1609\/icwsm.v16i1.19330"},{"key":"e_1_3_2_2_40_1","doi-asserted-by":"publisher","DOI":"10.1609\/icwsm.v14i1.7354"},{"key":"e_1_3_2_2_41_1","first-page":"372","volume-title":"Ananthram Swami. The Limitations of Deep Learning in Adversarial Settings. In IEEE European Symposium on Security and Privacy (Euro S&P)","author":"Papernot Nicolas","year":"2016","unstructured":"Nicolas Papernot , Patrick D. McDaniel , Somesh Jha , Matt Fredrikson , Z. Berkay Celik , and Ananthram Swami. The Limitations of Deep Learning in Adversarial Settings. In IEEE European Symposium on Security and Privacy (Euro S&P) , pages 372 -- 387 . IEEE, 2016 . Nicolas Papernot, Patrick D. McDaniel, Somesh Jha, Matt Fredrikson, Z. Berkay Celik, and Ananthram Swami. The Limitations of Deep Learning in Adversarial Settings. In IEEE European Symposium on Security and Privacy (Euro S&P), pages 372--387. IEEE, 2016."},{"key":"e_1_3_2_2_42_1","doi-asserted-by":"publisher","DOI":"10.1109\/SP.2016.41"},{"key":"e_1_3_2_2_43_1","first-page":"311","volume-title":"Annual Meeting of the Association for Computational Linguistics (ACL)","author":"Papineni Kishore","year":"2016","unstructured":"Kishore Papineni , Salim Roukos , Todd Ward , and Wei-Jing Zhu . Bleu : a Method for Automatic Evaluation of Machine Translation . In Annual Meeting of the Association for Computational Linguistics (ACL) , pages 311 -- 318 . ACL, 2016 . Kishore Papineni, Salim Roukos, Todd Ward, and Wei-Jing Zhu. Bleu: a Method for Automatic Evaluation of Machine Translation. In Annual Meeting of the Association for Computational Linguistics (ACL), pages 311--318. ACL, 2016."},{"key":"e_1_3_2_2_44_1","volume-title":"Red Teaming Language Models with Language Models. CoRR abs\/2202.03286","author":"Perez Ethan","year":"2022","unstructured":"Ethan Perez , Saffron Huang , H. Francis Song , Trevor Cai , Roman Ring , John Aslanides , Amelia Glaese , Nat McAleese , and Geoffrey Irving . Red Teaming Language Models with Language Models. CoRR abs\/2202.03286 , 2022 . Ethan Perez, Saffron Huang, H. Francis Song, Trevor Cai, Roman Ring, John Aslanides, Amelia Glaese, Nat McAleese, and Geoffrey Irving. Red Teaming Language Models with Language Models. CoRR abs\/2202.03286, 2022."},{"key":"e_1_3_2_2_45_1","volume-title":"Language Models are Unsupervised Multitask Learners. OpenAI blog","author":"Radford Alec","year":"2019","unstructured":"Alec Radford , Jeffrey Wu , Rewon Child , David Luan , Dario Amodei , and Ilya Sutskever . Language Models are Unsupervised Multitask Learners. OpenAI blog , 2019 . Alec Radford, Jeffrey Wu, Rewon Child, David Luan, Dario Amodei, and Ilya Sutskever. Language Models are Unsupervised Multitask Learners. OpenAI blog, 2019."},{"key":"e_1_3_2_2_46_1","first-page":"557","volume-title":"Community-Specific Learning: How Distinctive Toxicity Norms Are Maintained in Political Subreddits. In International Conference on Web and Social Media (ICWSM)","author":"Rajadesingan Ashwin","year":"2020","unstructured":"Ashwin Rajadesingan , Paul Resnick , and Ceren Budak . Quick , Community-Specific Learning: How Distinctive Toxicity Norms Are Maintained in Political Subreddits. In International Conference on Web and Social Media (ICWSM) , pages 557 -- 568 . AAAI, 2020 . Ashwin Rajadesingan, Paul Resnick, and Ceren Budak. Quick, Community-Specific Learning: How Distinctive Toxicity Norms Are Maintained in Political Subreddits. In International Conference on Web and Social Media (ICWSM), pages 557--568. AAAI, 2020."},{"key":"e_1_3_2_2_47_1","first-page":"196","volume-title":"International Conference on Web and Social Media (ICWSM)","author":"Ribeiro Manoel Horta","year":"2021","unstructured":"Manoel Horta Ribeiro , Jeremy Blackburn , Barry Bradlyn , Emiliano De Cristofaro , Gianluca Stringhini , Summer Long , Stephanie Greenberg , and Savvas Zannettou . The Evolution of the Manosphere across the Web . In International Conference on Web and Social Media (ICWSM) , pages 196 -- 207 . AAAI, 2021 . Manoel Horta Ribeiro, Jeremy Blackburn, Barry Bradlyn, Emiliano De Cristofaro, Gianluca Stringhini, Summer Long, Stephanie Greenberg, and Savvas Zannettou. The Evolution of the Manosphere across the Web. In International Conference on Web and Social Media (ICWSM), pages 196--207. AAAI, 2021."},{"key":"e_1_3_2_2_48_1","first-page":"172","volume-title":"Bill Dolan. Unsupervised Modeling of Twitter Conversations. In Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (NAACL-HLT)","author":"Ritter Alan","year":"2010","unstructured":"Alan Ritter , Colin Cherry , and Bill Dolan. Unsupervised Modeling of Twitter Conversations. In Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (NAACL-HLT) , pages 172 -- 180 . ACL, 2010 . Alan Ritter, Colin Cherry, and Bill Dolan. Unsupervised Modeling of Twitter Conversations. In Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (NAACL-HLT), pages 172--180. ACL, 2010."},{"key":"e_1_3_2_2_49_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2021.eacl-main.24"},{"key":"e_1_3_2_2_50_1","volume-title":"Leveraging Pre-trained Checkpoints for Sequence Generation Tasks. CoRR abs\/1907.12461","author":"Rothe Sascha","year":"2019","unstructured":"Sascha Rothe , Shashi Narayan , and Aliaksei Severyn . Leveraging Pre-trained Checkpoints for Sequence Generation Tasks. CoRR abs\/1907.12461 , 2019 . Sascha Rothe, Shashi Narayan, and Aliaksei Severyn. Leveraging Pre-trained Checkpoints for Sequence Generation Tasks. CoRR abs\/1907.12461, 2019."},{"key":"e_1_3_2_2_51_1","first-page":"1668","volume-title":"Smith. The Risk of Racial Bias in Hate Speech Detection. In Annual Meeting of the Association for Computational Linguistics (ACL)","author":"Sap Maarten","year":"2019","unstructured":"Maarten Sap , Dallas Card , Saadia Gabriel , Yejin Choi , and Noah A . Smith. The Risk of Racial Bias in Hate Speech Detection. In Annual Meeting of the Association for Computational Linguistics (ACL) , pages 1668 -- 1678 . ACL, 2019 . Maarten Sap, Dallas Card, Saadia Gabriel, Yejin Choi, and Noah A. Smith. The Risk of Racial Bias in Hate Speech Detection. In Annual Meeting of the Association for Computational Linguistics (ACL), pages 1668--1678. ACL, 2019."},{"key":"e_1_3_2_2_52_1","first-page":"5248","volume-title":"Dirk Hovy. Predictive Biases in Natural Language Processing Models: A Conceptual Framework and Overview. In Annual Meeting of the Association for Computational Linguistics (ACL)","author":"Shah Deven","year":"2020","unstructured":"Deven Shah , H. Andrew Schwartz , and Dirk Hovy. Predictive Biases in Natural Language Processing Models: A Conceptual Framework and Overview. In Annual Meeting of the Association for Computational Linguistics (ACL) , pages 5248 -- 5264 . ACL, 2020 . Deven Shah, H. Andrew Schwartz, and Dirk Hovy. Predictive Biases in Natural Language Processing Models: A Conceptual Framework and Overview. In Annual Meeting of the Association for Computational Linguistics (ACL), pages 5248--5264. ACL, 2020."},{"key":"e_1_3_2_2_53_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/D19-1339"},{"key":"e_1_3_2_2_54_1","volume-title":"Towards Controllable Biases in Language Generation. CoRR abs\/2005.00268","author":"Sheng Emily","year":"2020","unstructured":"Emily Sheng , Kai-Wei Chang , Premkumar Natarajan , and Nanyun Peng . Towards Controllable Biases in Language Generation. CoRR abs\/2005.00268 , 2020 . Emily Sheng, Kai-Wei Chang, Premkumar Natarajan, and Nanyun Peng. Towards Controllable Biases in Language Generation. CoRR abs\/2005.00268, 2020."},{"key":"e_1_3_2_2_55_1","first-page":"3","volume-title":"Vitaly Shmatikov. Membership Inference Attacks Against Machine Learning Models. In IEEE Symposium on Security and Privacy (S&P)","author":"Shokri Reza","year":"2017","unstructured":"Reza Shokri , Marco Stronati , Congzheng Song , and Vitaly Shmatikov. Membership Inference Attacks Against Machine Learning Models. In IEEE Symposium on Security and Privacy (S&P) , pages 3 -- 18 . IEEE, 2017 . Reza Shokri, Marco Stronati, Congzheng Song, and Vitaly Shmatikov. Membership Inference Attacks Against Machine Learning Models. In IEEE Symposium on Security and Privacy (S&P), pages 3--18. IEEE, 2017."},{"key":"e_1_3_2_2_56_1","first-page":"2453","volume-title":"The Dialogue Dodecathlon: Open-Domain Knowledge and Image Grounded Conversational Agents. In Annual Meeting of the Association for Computational Linguistics (ACL)","author":"Shuster Kurt","year":"2020","unstructured":"Kurt Shuster , Da Ju , Stephen Roller , Emily Dinan , Y- Lan Boureau , and Jason We- ston. The Dialogue Dodecathlon: Open-Domain Knowledge and Image Grounded Conversational Agents. In Annual Meeting of the Association for Computational Linguistics (ACL) , pages 2453 -- 2470 . ACL, 2020 . Kurt Shuster, Da Ju, Stephen Roller, Emily Dinan, Y-Lan Boureau, and Jason We- ston. The Dialogue Dodecathlon: Open-Domain Knowledge and Image Grounded Conversational Agents. In Annual Meeting of the Association for Computational Linguistics (ACL), pages 2453--2470. ACL, 2020."},{"key":"e_1_3_2_2_57_1","first-page":"377","volume-title":"Song and Ananth Raghunathan. Information Leakage in Embedding Models. In ACM SIGSAC Conference on Computer and Communications Security (CCS)","author":"Congzheng","year":"2020","unstructured":"Congzheng Song and Ananth Raghunathan. Information Leakage in Embedding Models. In ACM SIGSAC Conference on Computer and Communications Security (CCS) , pages 377 -- 390 . ACM, 2020 . Congzheng Song and Ananth Raghunathan. Information Leakage in Embedding Models. In ACM SIGSAC Conference on Computer and Communications Security (CCS), pages 377--390. ACM, 2020."},{"key":"e_1_3_2_2_58_1","volume-title":"On the Safety of Conversational Models: Taxonomy, Dataset, and Benchmark. CoRR abs\/2110.08466","author":"Sun Hao","year":"2021","unstructured":"Hao Sun , Guangxuan Xu , Jiawen Deng , Jiale Cheng , Chujie Zheng , Hao Zhou , Nanyun Peng , Xiaoyan Zhu , and Minlie Huang . On the Safety of Conversational Models: Taxonomy, Dataset, and Benchmark. CoRR abs\/2110.08466 , 2021 . Hao Sun, Guangxuan Xu, Jiawen Deng, Jiale Cheng, Chujie Zheng, Hao Zhou, Nanyun Peng, Xiaoyan Zhu, and Minlie Huang. On the Safety of Conversational Models: Taxonomy, Dataset, and Benchmark. CoRR abs\/2110.08466, 2021."},{"key":"e_1_3_2_2_59_1","first-page":"3104","volume-title":"Annual Conference on Neural Information Processing Systems (NIPS)","author":"Sutskever Ilya","year":"2014","unstructured":"Ilya Sutskever , Oriol Vinyals , and Quoc V. Le . Sequence to Sequence Learning with Neural Networks . In Annual Conference on Neural Information Processing Systems (NIPS) , pages 3104 -- 3112 . NIPS, 2014 . Ilya Sutskever, Oriol Vinyals, and Quoc V. Le. Sequence to Sequence Learning with Neural Networks. In Annual Conference on Neural Information Processing Systems (NIPS), pages 3104--3112. NIPS, 2014."},{"key":"e_1_3_2_2_60_1","doi-asserted-by":"publisher","DOI":"10.1145\/3442381.3450024"},{"key":"e_1_3_2_2_61_1","first-page":"601","volume-title":"USENIX Security Symposium (USENIX Security)","author":"Tram\u00e8r Florian","year":"2016","unstructured":"Florian Tram\u00e8r , Fan Zhang , Ari Juels , Michael K. Reiter , and Thomas Ristenpart . Stealing Machine Learning Models via Prediction APIs . In USENIX Security Symposium (USENIX Security) , pages 601 -- 618 . USENIX, 2016 . Florian Tram\u00e8r, Fan Zhang, Ari Juels, Michael K. Reiter, and Thomas Ristenpart. Stealing Machine Learning Models via Prediction APIs. In USENIX Security Symposium (USENIX Security), pages 601--618. USENIX, 2016."},{"key":"e_1_3_2_2_62_1","volume-title":"SaFeRDialogues: Taking Feedback Gracefully after Conversational Safety Failures. CoRR abs\/2110.07518","author":"Ung Megan","year":"2021","unstructured":"Megan Ung , Jing Xu , and Y- Lan Boureau . SaFeRDialogues: Taking Feedback Gracefully after Conversational Safety Failures. CoRR abs\/2110.07518 , 2021 . Megan Ung, Jing Xu, and Y-Lan Boureau. SaFeRDialogues: Taking Feedback Gracefully after Conversational Safety Failures. CoRR abs\/2110.07518, 2021."},{"key":"e_1_3_2_2_63_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/D19-1062"},{"key":"e_1_3_2_2_64_1","first-page":"5998","volume-title":"Annual Conference on Neural Information Processing Systems (NIPS)","author":"Vaswani Ashish","year":"2017","unstructured":"Ashish Vaswani , Noam Shazeer , Niki Parmar , Jakob Uszkoreit , Llion Jones , Aidan N. Gomez , Lukasz Kaiser , and Illia Polosukhin . Attention is All you Need . In Annual Conference on Neural Information Processing Systems (NIPS) , pages 5998 -- 6008 . NIPS, 2017 . Ashish Vaswani, Noam Shazeer, Niki Parmar, Jakob Uszkoreit, Llion Jones, Aidan N. Gomez, Lukasz Kaiser, and Illia Polosukhin. Attention is All you Need. In Annual Conference on Neural Information Processing Systems (NIPS), pages 5998--6008. NIPS, 2017."},{"key":"e_1_3_2_2_65_1","first-page":"7371","volume-title":"Dhruv Batra. Diverse Beam Search for Improved Description of Complex Scenes. In AAAI Conference on Artificial Intelligence (AAAI)","author":"Vijayakumar Ashwin K.","year":"2018","unstructured":"Ashwin K. Vijayakumar , Michael Cogswell , Ramprasaath R. Selvaraju , Qing Sun , Stefan Lee , David J. Crandall , and Dhruv Batra. Diverse Beam Search for Improved Description of Complex Scenes. In AAAI Conference on Artificial Intelligence (AAAI) , pages 7371 -- 7379 . AAAI, 2018 . Ashwin K. Vijayakumar, Michael Cogswell, Ramprasaath R. Selvaraju, Qing Sun, Stefan Lee, David J. Crandall, and Dhruv Batra. Diverse Beam Search for Improved Description of Complex Scenes. In AAAI Conference on Artificial Intelligence (AAAI), pages 7371--7379. AAAI, 2018."},{"key":"e_1_3_2_2_66_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/D19-1221"},{"key":"e_1_3_2_2_67_1","first-page":"33","volume-title":"Qun Liu. Semantics-Enhanced Task-Oriented Dialogue Translation: A Case Study on Hotel Booking. In International Joint Conference on Natural Language Processing (IJCNLP)","author":"Wang Longyue","year":"2017","unstructured":"Longyue Wang , Jinhua Du , Liangyou Li , Zhaopeng Tu , Andy Way , and Qun Liu. Semantics-Enhanced Task-Oriented Dialogue Translation: A Case Study on Hotel Booking. In International Joint Conference on Natural Language Processing (IJCNLP) , pages 33 -- 36 . ACL, 2017 . Longyue Wang, Jinhua Du, Liangyou Li, Zhaopeng Tu, Andy Way, and Qun Liu. Semantics-Enhanced Task-Oriented Dialogue Translation: A Case Study on Hotel Booking. In International Joint Conference on Natural Language Processing (IJCNLP), pages 33--36. ACL, 2017."},{"key":"e_1_3_2_2_68_1","volume-title":"Ming Zhou. MiniLM: Deep Self-Attention Distillation for Task-Agnostic Compression of Pre- Trained Transformers. In Annual Conference on Neural Information Processing Systems (NeurIPS). NeurIPS","author":"Wang Wenhui","year":"2020","unstructured":"Wenhui Wang , Furu Wei , Li Dong , Hangbo Bao , Nan Yang , and Ming Zhou. MiniLM: Deep Self-Attention Distillation for Task-Agnostic Compression of Pre- Trained Transformers. In Annual Conference on Neural Information Processing Systems (NeurIPS). NeurIPS , 2020 . Wenhui Wang, Furu Wei, Li Dong, Hangbo Bao, Nan Yang, and Ming Zhou. MiniLM: Deep Self-Attention Distillation for Task-Agnostic Compression of Pre- Trained Transformers. In Annual Conference on Neural Information Processing Systems (NeurIPS). NeurIPS, 2020."},{"key":"e_1_3_2_2_69_1","volume-title":"Kirsty Anderson, Pushmeet Kohli, Ben Coppin, and Po-Sen Huang. Challenges in Detoxifying Language Models. CoRR abs\/2109.07445","author":"Welbl Johannes","year":"2021","unstructured":"Johannes Welbl , Amelia Glaese , Jonathan Uesato , Sumanth Dathathri , John Mellor , Lisa Anne Hendricks , Kirsty Anderson, Pushmeet Kohli, Ben Coppin, and Po-Sen Huang. Challenges in Detoxifying Language Models. CoRR abs\/2109.07445 , 2021 . Johannes Welbl, Amelia Glaese, Jonathan Uesato, Sumanth Dathathri, John Mellor, Lisa Anne Hendricks, Kirsty Anderson, Pushmeet Kohli, Ben Coppin, and Po-Sen Huang. Challenges in Detoxifying Language Models. CoRR abs\/2109.07445, 2021."},{"key":"e_1_3_2_2_70_1","first-page":"11530","volume-title":"AAAI Conference on Artificial Intelligence (AAAI)","author":"Xu Canwen","year":"2022","unstructured":"Canwen Xu , Zexue He , Zhankui He , and Julian J . McAuley. Leashing the In- ner Demons: Self-Detoxification for Language Models . In AAAI Conference on Artificial Intelligence (AAAI) , pages 11530 -- 11537 . AAAI, 2022 . Canwen Xu, Zexue He, Zhankui He, and Julian J. McAuley. Leashing the In- ner Demons: Self-Detoxification for Language Models. In AAAI Conference on Artificial Intelligence (AAAI), pages 11530--11537. AAAI, 2022."},{"key":"e_1_3_2_2_71_1","volume-title":"Recipes for Safety in Open-domain Chatbots. CoRR abs\/2010.07079","author":"Xu Jing","year":"2020","unstructured":"Jing Xu , Da Ju , Margaret Li , Y- Lan Boureau , Jason Weston , and Emily Dinan . Recipes for Safety in Open-domain Chatbots. CoRR abs\/2010.07079 , 2020 . Jing Xu, Da Ju, Margaret Li, Y-Lan Boureau, Jason Weston, and Emily Dinan. Recipes for Safety in Open-domain Chatbots. CoRR abs\/2010.07079, 2020."},{"key":"e_1_3_2_2_72_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/2021.naacl-main.235"},{"key":"e_1_3_2_2_73_1","doi-asserted-by":"publisher","DOI":"10.1609\/aaai.v31i1.11182"},{"key":"e_1_3_2_2_74_1","doi-asserted-by":"publisher","DOI":"10.1145\/3278532.3278550"},{"key":"e_1_3_2_2_75_1","first-page":"405","volume-title":"Jeremy Blackburn. The Web Centipede: Understanding How Web Communities Influence Each Other Through the Lens of Mainstream and Alternative News Sources. In ACM Internet Measurement Conference (IMC)","author":"Zannettou Savvas","year":"2017","unstructured":"Savvas Zannettou , Tristan Caulfield , Emiliano De Cristofaro , Nicolas Kourtellis , Ilias Leontiadis , Michael Sirivianos , Gianluca Stringhini , and Jeremy Blackburn. The Web Centipede: Understanding How Web Communities Influence Each Other Through the Lens of Mainstream and Alternative News Sources. In ACM Internet Measurement Conference (IMC) , pages 405 -- 417 . ACM, 2017 . Savvas Zannettou, Tristan Caulfield, Emiliano De Cristofaro, Nicolas Kourtellis, Ilias Leontiadis, Michael Sirivianos, Gianluca Stringhini, and Jeremy Blackburn. The Web Centipede: Understanding How Web Communities Influence Each Other Through the Lens of Mainstream and Alternative News Sources. In ACM Internet Measurement Conference (IMC), pages 405--417. ACM, 2017."},{"key":"e_1_3_2_2_76_1","first-page":"125","volume-title":"Stringhini. Measuring and Characterizing Hate Speech on News Websites. In ACM Conference on Web Science (WebSci)","author":"Zannettou Savvas","year":"2020","unstructured":"Savvas Zannettou , Mai ElSherief , Elizabeth M. Belding , Shirin Nilizadeh , and Gi- anluca Stringhini. Measuring and Characterizing Hate Speech on News Websites. In ACM Conference on Web Science (WebSci) , pages 125 -- 134 . ACM, 2020 . Savvas Zannettou, Mai ElSherief, Elizabeth M. Belding, Shirin Nilizadeh, and Gi- anluca Stringhini. Measuring and Characterizing Hate Speech on News Websites. In ACM Conference on Web Science (WebSci), pages 125--134. ACM, 2020."},{"key":"e_1_3_2_2_77_1","doi-asserted-by":"publisher","DOI":"10.1609\/icwsm.v14i1.7343"},{"key":"e_1_3_2_2_78_1","volume-title":"DialoGPT: Large-Scale Generative Pre-training for Conversational Response Generation. CoRR abs\/1911.00536","author":"Zhang Yizhe","year":"2019","unstructured":"Yizhe Zhang , Siqi Sun , Michel Galley , Yen-Chun Chen , Chris Brockett , Xiang Gao , Jianfeng Gao , Jingjing Liu , and Bill Dolan . DialoGPT: Large-Scale Generative Pre-training for Conversational Response Generation. CoRR abs\/1911.00536 , 2019 . Yizhe Zhang, Siqi Sun, Michel Galley, Yen-Chun Chen, Chris Brockett, Xiang Gao, Jianfeng Gao, Jingjing Liu, and Bill Dolan. DialoGPT: Large-Scale Generative Pre-training for Conversational Response Generation. CoRR abs\/1911.00536, 2019."},{"key":"e_1_3_2_2_79_1","first-page":"1097","volume-title":"Yong Yu. Texygen: A Benchmarking Platform for Text Generation Models. In International ACM SIGIR Conference on Research and Development in Information Retrieval (SIGIR)","author":"Zhu Yaoming","year":"2018","unstructured":"Yaoming Zhu , Sidi Lu , Lei Zheng , Jiaxian Guo , Weinan Zhang , Jun Wang , and Yong Yu. Texygen: A Benchmarking Platform for Text Generation Models. In International ACM SIGIR Conference on Research and Development in Information Retrieval (SIGIR) , pages 1097 -- 1100 . ACM, 2018 . Yaoming Zhu, Sidi Lu, Lei Zheng, Jiaxian Guo, Weinan Zhang, Jun Wang, and Yong Yu. Texygen: A Benchmarking Platform for Text Generation Models. In International ACM SIGIR Conference on Research and Development in Information Retrieval (SIGIR), pages 1097--1100. ACM, 2018."},{"key":"e_1_3_2_2_80_1","volume-title":"Racism is a Virus: Anti- Asian Hate and Counterhate in Social Media during the COVID-19 Crisis. CoRR abs\/2005.12423","author":"Ziems Caleb","year":"2021","unstructured":"Caleb Ziems , Bing He , Sandeep Soni , and Srijan Kumar . Racism is a Virus: Anti- Asian Hate and Counterhate in Social Media during the COVID-19 Crisis. CoRR abs\/2005.12423 , 2021 . Caleb Ziems, Bing He, Sandeep Soni, and Srijan Kumar. Racism is a Virus: Anti- Asian Hate and Counterhate in Social Media during the COVID-19 Crisis. CoRR abs\/2005.12423, 2021."}],"event":{"name":"CCS '22: 2022 ACM SIGSAC Conference on Computer and Communications Security","location":"Los Angeles CA USA","acronym":"CCS '22","sponsor":["SIGSAC ACM Special Interest Group on Security, Audit, and Control"]},"container-title":["Proceedings of the 2022 ACM SIGSAC Conference on Computer and Communications Security"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3548606.3560599","content-type":"unspecified","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3548606.3560599","content-type":"application\/pdf","content-version":"vor","intended-application":"syndication"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3548606.3560599","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,6,17]],"date-time":"2025-06-17T17:50:58Z","timestamp":1750182658000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3548606.3560599"}},"subtitle":["Measuring and Triggering Toxic Behavior in Open-Domain Chatbots"],"short-title":[],"issued":{"date-parts":[[2022,11,7]]},"references-count":80,"alternative-id":["10.1145\/3548606.3560599","10.1145\/3548606"],"URL":"https:\/\/doi.org\/10.1145\/3548606.3560599","relation":{},"subject":[],"published":{"date-parts":[[2022,11,7]]},"assertion":[{"value":"2022-11-07","order":2,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}