{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,8,22]],"date-time":"2025-08-22T20:10:05Z","timestamp":1755893405455,"version":"3.44.0"},"publisher-location":"New York, NY, USA","reference-count":42,"publisher":"ACM","license":[{"start":{"date-parts":[[2023,12,15]],"date-time":"2023-12-15T00:00:00Z","timestamp":1702598400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0\/"}],"funder":[{"name":"Institute for Information & communications Technology Promotion","award":["2018-0-00749"],"award-info":[{"award-number":["2018-0-00749"]}]},{"DOI":"10.13039\/501100006374","name":"National Research Foundation of Korea","doi-asserted-by":"publisher","award":["2022R1A2C1012633"],"award-info":[{"award-number":["2022R1A2C1012633"]}],"id":[{"id":"10.13039\/501100006374","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100006374","name":"Korea International Cooperation Agency","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100006374","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2023,12,15]]},"DOI":"10.1145\/3639233.3639344","type":"proceedings-article","created":{"date-parts":[[2024,3,5]],"date-time":"2024-03-05T11:02:10Z","timestamp":1709636530000},"page":"63-70","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":0,"title":["Fine-Tuning BERT on Twitter and Reddit Data in Luganda and English"],"prefix":"10.1145","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-3027-4526","authenticated-orcid":false,"given":"Richard","family":"Kimera","sequence":"first","affiliation":[{"name":"Department of Advanced Convergence, Handong Global University, South Korea"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-4527-8028","authenticated-orcid":false,"given":"Daniela N","family":"Rim","sequence":"additional","affiliation":[{"name":"Department of Computer Science and Electrical Engineering, Handong Global University, South Korea"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0855-8725","authenticated-orcid":false,"given":"Heeyoul","family":"Choi","sequence":"additional","affiliation":[{"name":"Department of Computer Science and Electrical Engineering, Handong Global University, South Korea"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"320","published-online":{"date-parts":[[2024,3,5]]},"reference":[{"key":"e_1_3_2_1_1_1","doi-asserted-by":"publisher","DOI":"10.7763\/IJCTE.2020.V12.1280"},{"key":"e_1_3_2_1_2_1","doi-asserted-by":"publisher","DOI":"10.1162\/tacl_a_00288"},{"key":"e_1_3_2_1_3_1","first-page":"3","article-title":"Luganda Nouns","volume":"3","author":"Baertlein Elizabeth","year":"2014","unstructured":"Elizabeth Baertlein and Martin Ssekitto. 2014. Luganda Nouns: Inflectional Morphology and Tests. Linguistic Portfolios 3, 1 (2014), 3.","journal-title":"Inflectional Morphology and Tests. Linguistic Portfolios"},{"key":"e_1_3_2_1_4_1","volume-title":"Building machine translation systems for the next thousand languages. arXiv preprint arXiv:2205.03983","author":"Bapna Ankur","year":"2022","unstructured":"Ankur Bapna, Isaac Caswell, Julia Kreutzer, Orhan Firat, Daan van Esch, Aditya Siddhant, Mengmeng Niu, Pallavi Baljekar, Xavier Garcia, Wolfgang Macherey, 2022. Building machine translation systems for the next thousand languages. arXiv preprint arXiv:2205.03983 (2022)."},{"key":"e_1_3_2_1_5_1","volume-title":"Tweeteval: Unified benchmark and comparative evaluation for tweet classification. arXiv preprint arXiv:2010.12421","author":"Barbieri Francesco","year":"2020","unstructured":"Francesco Barbieri, Jose Camacho-Collados, Leonardo Neves, and Luis Espinosa-Anke. 2020. Tweeteval: Unified benchmark and comparative evaluation for tweet classification. arXiv preprint arXiv:2010.12421 (2020)."},{"key":"e_1_3_2_1_6_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/S18-1003"},{"key":"e_1_3_2_1_7_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/S19-2007"},{"key":"e_1_3_2_1_8_1","volume-title":"Byte pair encoding is suboptimal for language model pretraining. arXiv preprint arXiv:2004.03720","author":"Bostrom Kaj","year":"2020","unstructured":"Kaj Bostrom and Greg Durrett. 2020. Byte pair encoding is suboptimal for language model pretraining. arXiv preprint arXiv:2004.03720 (2020)."},{"key":"e_1_3_2_1_9_1","volume-title":"An Introduction to Survival Luganda. Kampala","author":"Byakutaga Shirley","year":"2008","unstructured":"Shirley Byakutaga, Henry Kabayo, Herbert Sengendo, and Ven Kitone. 2008. An Introduction to Survival Luganda. Kampala, Uganda: Peace Corps Uganda. Accessed July 11 (2008), 2017."},{"key":"e_1_3_2_1_10_1","unstructured":"Isaac Caswell. 2022. Google Translate learns 24 new languages. Google Blog. https:\/\/blog.google\/products\/translate\/24-new-languages\/Accessed on 2024\/01\/13 14:46:56."},{"key":"e_1_3_2_1_11_1","volume-title":"Rethinking embedding coupling in pre-trained language models. arXiv preprint arXiv:2010.12821","author":"Chung Hyung\u00a0Won","year":"2020","unstructured":"Hyung\u00a0Won Chung, Thibault Fevry, Henry Tsai, Melvin Johnson, and Sebastian Ruder. 2020. Rethinking embedding coupling in pre-trained language models. arXiv preprint arXiv:2010.12821 (2020)."},{"key":"e_1_3_2_1_12_1","volume-title":"Bert: Pre-training of deep bidirectional transformers for language understanding. arXiv preprint arXiv:1810.04805","author":"Devlin Jacob","year":"2018","unstructured":"Jacob Devlin, Ming-Wei Chang, Kenton Lee, and Kristina Toutanova. 2018. Bert: Pre-training of deep bidirectional transformers for language understanding. arXiv preprint arXiv:1810.04805 (2018)."},{"key":"e_1_3_2_1_13_1","volume-title":"Towards a common semantics for English count and mass nouns. Linguistics and philosophy 15","author":"Gillon S","year":"1992","unstructured":"Brendan\u00a0S Gillon. 1992. Towards a common semantics for English count and mass nouns. Linguistics and philosophy 15 (1992), 597\u2013639."},{"key":"e_1_3_2_1_14_1","volume-title":"Transfer Learning for Low-Resource Sentiment Analysis. arXiv preprint arXiv:2304.04703","author":"Hameed Razhan","year":"2023","unstructured":"Razhan Hameed, Sina Ahmadi, and Fatemeh Daneshfar. 2023. Transfer Learning for Low-Resource Sentiment Analysis. arXiv preprint arXiv:2304.04703 (2023)."},{"key":"e_1_3_2_1_15_1","article-title":"Text based personality prediction from multiple social media data sources using pre-trained language model and model averaging","volume":"8","author":"Hans C","year":"2021","unstructured":"C Hans, D Suhartono, C Andry, and KZ Zamli. 2021. Text based personality prediction from multiple social media data sources using pre-trained language model and model averaging. Journal of Big Data 8, 68 (2021).","journal-title":"Journal of Big Data"},{"key":"e_1_3_2_1_16_1","volume-title":"Debertav3: Improving deberta using electra-style pre-training with gradient-disentangled embedding sharing. arXiv preprint arXiv:2111.09543","author":"He Pengcheng","year":"2021","unstructured":"Pengcheng He, Jianfeng Gao, and Weizhu Chen. 2021. Debertav3: Improving deberta using electra-style pre-training with gradient-disentangled embedding sharing. arXiv preprint arXiv:2111.09543 (2021)."},{"key":"e_1_3_2_1_17_1","volume-title":"FakeBERT: Fake news detection in social media with a BERT-based deep learning approach. Multimedia tools and applications 80, 8","author":"Kaliyar Rohit\u00a0Kumar","year":"2021","unstructured":"Rohit\u00a0Kumar Kaliyar, Anurag Goswami, and Pratik Narang. 2021. FakeBERT: Fake news detection in social media with a BERT-based deep learning approach. Multimedia tools and applications 80, 8 (2021), 11765\u201311788."},{"key":"e_1_3_2_1_18_1","volume-title":"Building a Parallel Corpus and Training Translation Models Between Luganda and English. arXiv preprint arXiv:2301.02773","author":"Kimera Richard","year":"2023","unstructured":"Richard Kimera, Daniela\u00a0N Rim, and Heeyoul Choi. 2023. Building a Parallel Corpus and Training Translation Models Between Luganda and English. arXiv preprint arXiv:2301.02773 (2023)."},{"key":"e_1_3_2_1_19_1","volume-title":"Exploiting similarities among languages for machine translation. arXiv preprint arXiv:1309.4168","author":"Mikolov Tomas","year":"2013","unstructured":"Tomas Mikolov, Quoc\u00a0V Le, and Ilya Sutskever. 2013. Exploiting similarities among languages for machine translation. arXiv preprint arXiv:1309.4168 (2013)."},{"key":"e_1_3_2_1_20_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/S18-1001"},{"key":"e_1_3_2_1_21_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/S16-1003"},{"key":"e_1_3_2_1_22_1","doi-asserted-by":"publisher","DOI":"10.3390\/app11188575"},{"key":"e_1_3_2_1_23_1","volume-title":"Complex Networks and Their Applications VIII","author":"Mozafari Marzieh","year":"2019","unstructured":"Marzieh Mozafari, Reza Farahbakhsh, and Noel Crespi. 2020. A BERT-based transfer learning approach for hate speech detection in online social media. In Complex Networks and Their Applications VIII: Volume 1 Proceedings of the Eighth International Conference on Complex Networks and Their Applications COMPLEX NETWORKS 2019 8. Springer, 928\u2013940."},{"key":"e_1_3_2_1_24_1","volume-title":"Naijasenti: A nigerian twitter sentiment corpus for multilingual sentiment analysis. arXiv preprint arXiv:2201.08277","author":"Muhammad Shamsuddeen\u00a0Hassan","year":"2022","unstructured":"Shamsuddeen\u00a0Hassan Muhammad, David\u00a0Ifeoluwa Adelani, Sebastian Ruder, Ibrahim\u00a0Said Ahmad, Idris Abdulmumin, Bello\u00a0Shehu Bello, Monojit Choudhury, Chris\u00a0Chinenye Emezue, Saheed\u00a0Salahudeen Abdullahi, Anuoluwapo Aremu, 2022. Naijasenti: A nigerian twitter sentiment corpus for multilingual sentiment analysis. arXiv preprint arXiv:2201.08277 (2022)."},{"key":"e_1_3_2_1_25_1","first-page":"12","article-title":"Mining voter sentiments from Twitter data for the 2016 Uganda Presidential elections","volume":"3","author":"Mukonyezi Isaac","year":"2018","unstructured":"Isaac Mukonyezi, Claire Babirye, and Ernest Mwebaze. 2018. Mining voter sentiments from Twitter data for the 2016 Uganda Presidential elections. International Journal of Technology and Management 3, 2 (2018), 12\u201312.","journal-title":"International Journal of Technology and Management"},{"key":"e_1_3_2_1_26_1","doi-asserted-by":"publisher","DOI":"10.1093\/intqhc\/mzr040"},{"key":"e_1_3_2_1_27_1","volume-title":"Misinformation detection in Luganda-English code-mixed social media text. arXiv preprint arXiv:2104.00124","author":"Nabende Peter","year":"2021","unstructured":"Peter Nabende, David Kabiito, Claire Babirye, Hewitt Tusiime, and Joyce Nakatumba-Nabende. 2021. Misinformation detection in Luganda-English code-mixed social media text. arXiv preprint arXiv:2104.00124 (2021)."},{"key":"e_1_3_2_1_28_1","doi-asserted-by":"publisher","DOI":"10.1186\/s40649-019-0063-4"},{"key":"e_1_3_2_1_29_1","volume-title":"KinyaBERT: a morphology-aware Kinyarwanda language model. arXiv preprint arXiv:2203.08459","author":"Nzeyimana Antoine","year":"2022","unstructured":"Antoine Nzeyimana and Andre\u00a0Niyongabo Rubungo. 2022. KinyaBERT: a morphology-aware Kinyarwanda language model. arXiv preprint arXiv:2203.08459 (2022)."},{"volume-title":"AfriBERTa: Towards Viable Multilingual Language Models for Low-resource Languages. Master\u2019s thesis","author":"Ogueji Kelechi","key":"e_1_3_2_1_30_1","unstructured":"Kelechi Ogueji. 2022. AfriBERTa: Towards Viable Multilingual Language Models for Low-resource Languages. Master\u2019s thesis. University of Waterloo."},{"key":"e_1_3_2_1_31_1","doi-asserted-by":"publisher","DOI":"10.7763\/IJCTE.2021.V13.1297"},{"key":"e_1_3_2_1_32_1","doi-asserted-by":"publisher","DOI":"10.7763\/IJCTE.2021.V13.1287"},{"key":"e_1_3_2_1_33_1","volume-title":"SemEval-2017 task 4: Sentiment analysis in Twitter. arXiv preprint arXiv:1912.00741","author":"Rosenthal Sara","year":"2019","unstructured":"Sara Rosenthal, Noura Farra, and Preslav Nakov. 2019. SemEval-2017 task 4: Sentiment analysis in Twitter. arXiv preprint arXiv:1912.00741 (2019)."},{"key":"e_1_3_2_1_34_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICDE.2013.6544931"},{"key":"e_1_3_2_1_35_1","volume-title":"Neural machine translation of rare words with subword units. arXiv preprint arXiv:1508.07909","author":"Sennrich Rico","year":"2015","unstructured":"Rico Sennrich, Barry Haddow, and Alexandra Birch. 2015. Neural machine translation of rare words with subword units. arXiv preprint arXiv:1508.07909 (2015)."},{"key":"e_1_3_2_1_36_1","volume-title":"Enhancing African low-resource languages: Swahili data for language modelling. Data in brief 31","author":"Shikali S","year":"2020","unstructured":"Casper\u00a0S Shikali and Refuoe Mokhosi. 2020. Enhancing African low-resource languages: Swahili data for language modelling. Data in brief 31 (2020), 105951."},{"key":"e_1_3_2_1_37_1","doi-asserted-by":"publisher","DOI":"10.9728\/dcs.2020.21.5.951"},{"key":"e_1_3_2_1_38_1","unstructured":"Uganda National Health Users\u2019\/Consumers\u2019 Organisation (UNHCO). 2003. Study on Patient Feedback Mechanisms at Health Facilities in Uganda. Technical Report. Uganda National Health Users\u2019\/Consumers\u2019 Organisation (UNHCO). http:\/\/unhco.or.ug\/wp-content\/uploads\/downloads\/2010\/12\/UNHCO-Report-on-Feedback-Mechanisms-at-Health-Facilities-2003.pdf"},{"key":"e_1_3_2_1_39_1","doi-asserted-by":"publisher","DOI":"10.18653\/v1\/S18-1005"},{"key":"e_1_3_2_1_40_1","volume-title":"Attention is all you need. Advances in neural information processing systems 30","author":"Vaswani Ashish","year":"2017","unstructured":"Ashish Vaswani, Noam Shazeer, Niki Parmar, Jakob Uszkoreit, Llion Jones, Aidan\u00a0N Gomez, \u0141ukasz Kaiser, and Illia Polosukhin. 2017. Attention is all you need. Advances in neural information processing systems 30 (2017)."},{"key":"e_1_3_2_1_41_1","volume-title":"Has Sentiment Returned to the Pre-pandemic Level? A Sentiment Analysis Using US College Subreddit Data from 2019 to","author":"Yan Tian","year":"2022","unstructured":"Tian Yan and Fang Liu. 2023. Has Sentiment Returned to the Pre-pandemic Level? A Sentiment Analysis Using US College Subreddit Data from 2019 to 2022. arXiv preprint arXiv:2309.08845 (2023)."},{"key":"e_1_3_2_1_42_1","volume-title":"Semeval-2019 task 6: Identifying and categorizing offensive language in social media (offenseval). arXiv preprint arXiv:1903.08983","author":"Zampieri Marcos","year":"2019","unstructured":"Marcos Zampieri, Shervin Malmasi, Preslav Nakov, Sara Rosenthal, Noura Farra, and Ritesh Kumar. 2019. Semeval-2019 task 6: Identifying and categorizing offensive language in social media (offenseval). arXiv preprint arXiv:1903.08983 (2019)."}],"event":{"name":"NLPIR 2023: 2023 7th International Conference on Natural Language Processing and Information Retrieval","acronym":"NLPIR 2023","location":"Seoul Republic of Korea"},"container-title":["Proceedings of the 2023 7th International Conference on Natural Language Processing and Information Retrieval"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3639233.3639344","content-type":"unspecified","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/3639233.3639344","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,8,22]],"date-time":"2025-08-22T19:57:18Z","timestamp":1755892638000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/3639233.3639344"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,12,15]]},"references-count":42,"alternative-id":["10.1145\/3639233.3639344","10.1145\/3639233"],"URL":"https:\/\/doi.org\/10.1145\/3639233.3639344","relation":{},"subject":[],"published":{"date-parts":[[2023,12,15]]},"assertion":[{"value":"2024-03-05","order":3,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}