{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,1,12]],"date-time":"2026-01-12T23:12:03Z","timestamp":1768259523101,"version":"3.49.0"},"reference-count":26,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2025,1,9]],"date-time":"2025-01-09T00:00:00Z","timestamp":1736380800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,1,9]],"date-time":"2025-01-09T00:00:00Z","timestamp":1736380800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Lang Resources &amp; Evaluation"],"published-print":{"date-parts":[[2025,9]]},"DOI":"10.1007\/s10579-024-09803-2","type":"journal-article","created":{"date-parts":[[2025,1,9]],"date-time":"2025-01-09T08:11:48Z","timestamp":1736410308000},"page":"2169-2184","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Introducing a Swahili social media sentiment analysis dataset for the telecom industry"],"prefix":"10.1007","volume":"59","author":[{"given":"Mahadia","family":"Tunga","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Davis","family":"David","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,1,9]]},"reference":[{"key":"9803_CR1","unstructured":"Alabi, O. J., Adelani, I. D., Mosbach, M., & Klakow. D. (2022). Adapting pretrained language models to African languages via multilingual adaptive fine-tuning. https:\/\/aclanthology.org\/2022.acl-long.265\/"},{"key":"9803_CR2","doi-asserted-by":"publisher","unstructured":"Alshaabi, T., Dewhurst D. R., Minot J. R., Arnold M. V., Adams J. L., Danforth C. M., & Dodds P. S. (2021).The growing amplification of social media: Measuring temporal and social contagion dynamics for over 150 languages on Twitter for 2009\u20132020. https:\/\/doi.org\/10.48550\/arXiv.2003.03667","DOI":"10.48550\/arXiv.2003.03667"},{"key":"9803_CR3","doi-asserted-by":"publisher","DOI":"10.48550\/arXiv.2105.14373","author":"S Barreto","year":"2021","unstructured":"Barreto, S., Moura, R., Carvalho, J., Paes, A., & Plastino, A. (2021). Sentiment analysis in tweets: an assessment study from classical to modern text representation models. Data Mining and Knowledge Discovery. https:\/\/doi.org\/10.48550\/arXiv.2105.14373","journal-title":"Data Mining and Knowledge Discovery"},{"key":"9803_CR4","doi-asserted-by":"publisher","unstructured":"Conneau, A., Khandelwal, K., Goyal, N., Chaudhary, V., Wenzek, G., Guzm\u00e1n, F., Grave, E., Ott, M., Zettlemoyer, L., & Veselin Stoyanov, V. (2020). Unsupervised cross-lingual representation learning at scale. https:\/\/doi.org\/10.48550\/arXiv.1911.02116","DOI":"10.48550\/arXiv.1911.02116"},{"key":"9803_CR5","doi-asserted-by":"publisher","unstructured":"Das, B., & Chakraborty, S. (2018). An improved text sentiment classification model using TF-IDF and next word negation. https:\/\/doi.org\/10.48550\/arXiv.1806.06407","DOI":"10.48550\/arXiv.1806.06407"},{"key":"9803_CR6","doi-asserted-by":"publisher","DOI":"10.48550\/arXiv.2110.05719","author":"MA Davani","year":"2021","unstructured":"Davani, M. A., D\u2019iaz, M., & Prabhakaran, V. (2021). Dealing with disagreements: Looking beyond the majority vote in subjective annotations. Transactions of the Association. https:\/\/doi.org\/10.48550\/arXiv.2110.05719","journal-title":"Transactions of the Association"},{"key":"9803_CR7","doi-asserted-by":"publisher","unstructured":"Galke, L., & Scherp, A. (2021). Bag-of-words vs. graph vs. sequence in text classification: Questioning the necessity of text-graphs and the surprising strength of a wide MLP. https:\/\/doi.org\/10.48550\/arXiv.2109.03777","DOI":"10.48550\/arXiv.2109.03777"},{"key":"9803_CR8","doi-asserted-by":"publisher","DOI":"10.14569\/IJACSA.2020.0111190","author":"AA Ibrahim","year":"2020","unstructured":"Ibrahim, A. A., Ridwan, R. L., Muhammed, M. M., Abdulaziz, R. O., & Saheed, G. A. (2020). Comparison of the CatBoost classifier with other machine learning methods. International Journal of Advanced Computer Science and Applications (IJACSA). https:\/\/doi.org\/10.14569\/IJACSA.2020.0111190","journal-title":"International Journal of Advanced Computer Science and Applications (IJACSA)"},{"key":"9803_CR9","unstructured":"International Monetary Fund (IMF). (2021). World economic outlook database. Retrieved July 27, 2023, from https:\/\/www.imf.org\/en\/Publications\/WEO\/weo-database\/2021\/October"},{"key":"9803_CR10","doi-asserted-by":"publisher","unstructured":"Kapur, K., & Harikrishnan, R. (2022). Comparative study of sentiment analysis for multi-sourced social media platforms. https:\/\/doi.org\/10.48550\/arXiv.2212.04688","DOI":"10.48550\/arXiv.2212.04688"},{"key":"9803_CR11","doi-asserted-by":"publisher","DOI":"10.1017\/CBO9781139084789","volume-title":"Sentiment analysis: Mining opinions, sentiments, and emotions","author":"B Liu","year":"2020","unstructured":"Liu, B. (2020). Sentiment analysis: Mining opinions, sentiments, and emotions. Cambridge University Press. https:\/\/doi.org\/10.1017\/CBO9781139084789"},{"key":"9803_CR12","unstructured":"Liu, Q., Kusner J. M., & Blunsom, P. (2020). A survey on contextual embeddings. https:\/\/arxiv.org\/pdf\/2003.07278"},{"key":"9803_CR13","doi-asserted-by":"publisher","unstructured":"Mao, Y., Zhou, L., & Xiong, N. (2020). Identify influential nodes in online social network for brand communication. https:\/\/doi.org\/10.48550\/arXiv.2006.14104","DOI":"10.48550\/arXiv.2006.14104"},{"key":"9803_CR14","doi-asserted-by":"publisher","unstructured":"Martin, G. L., Mswahili, M. E., & Jeong, Y. (2021). Sentiment classification in Swahili language using multilingual bert. https:\/\/doi.org\/10.48550\/arXiv.2104.09006","DOI":"10.48550\/arXiv.2104.09006"},{"key":"9803_CR15","doi-asserted-by":"crossref","unstructured":"Martin, L, G., Mswahili, E. M., Jeong, Y., & Woo, J. (2022). SwahBERT: Language model of Swahili. https:\/\/aclanthology.org\/2022.naacl-main.23.pdf","DOI":"10.18653\/v1\/2022.naacl-main.23"},{"key":"9803_CR16","doi-asserted-by":"publisher","unstructured":"Mohammad, S. M. (2016). A practical guide to sentiment annotation: Challenges and solutions. https:\/\/doi.org\/10.18653\/v1\/W16-0429","DOI":"10.18653\/v1\/W16-0429"},{"key":"9803_CR17","doi-asserted-by":"publisher","unstructured":"Muhammad, S. H., Abdulmumin, I., Ayele, A. A., Ousidhoum, N., Adelani, D. I., Yimam, S. M., Ahmad, I. S., Beloucif, M., Mohammad, S. M., Ruder, S., Hourrane, O., Brazdil, P., Ali, F. D., David, D., Osei, S., Bello, B. S., Ibrahim, F., Gwadabe, T., Rutunda, S., \u2026, Steven, S. (2023). AfriSenti: A Twitter sentiment analysis benchmark for African languages. https:\/\/doi.org\/10.48550\/arXiv.2302.08956","DOI":"10.48550\/arXiv.2302.08956"},{"key":"9803_CR18","doi-asserted-by":"publisher","unstructured":"Muhammad, S.H., Adelani, D. I., Ruder, S., Ahmad, I. S., Abdulmumin, I., Bello, B. S., Choudhury, M., Emezue, C. C., Abdullahi, S. S., Aremu, A., Jeorge, A., & Brazdil, P. (2022). NaijaSenti: A Nigerian Twitter sentiment corpus for multilingual sentiment analysis. European Language Resources Association (ELRA). https:\/\/doi.org\/10.48550\/arXiv.2201.08277","DOI":"10.48550\/arXiv.2201.08277"},{"key":"9803_CR19","unstructured":"Muraiana, O. I. (2022). Ideal dataset splitting ratios in machine learning algorithms: General concerns for data scientists and data analysts. 7th International Mardin Artuklu Scientific Researches Conference."},{"key":"9803_CR20","doi-asserted-by":"crossref","unstructured":"Ogueji, K., Zhu, Y., & Lin, J. (2021). Small data? No problem! Exploring the viability of pretrained multilingual language models for low-resourced languages. https:\/\/aclanthology.org\/2021.mrl-1.11\/","DOI":"10.18653\/v1\/2021.mrl-1.11"},{"key":"9803_CR21","doi-asserted-by":"publisher","DOI":"10.3390\/info11060314","author":"J Samuel","year":"2020","unstructured":"Samuel, J., Ali, G. M. N., Rahman, M. M., Esawi, E., & Samuel, Y. (2020). Covid-19 public sentiment insights and machine learning for tweets classification. Information. https:\/\/doi.org\/10.3390\/info11060314","journal-title":"Information"},{"key":"9803_CR22","unstructured":"Shode, I., Adelani, I. D., & Anna Feldman, A. (2022). YOSM: A New Yoruba sentiment corpus for movie reviews. https:\/\/arxiv.org\/pdf\/2204.09711"},{"key":"9803_CR23","doi-asserted-by":"publisher","unstructured":"Singh, G. (2021). Sentiment analysis of code-mixed social media text (Hinglish). https:\/\/doi.org\/10.48550\/arXiv.2102.12149","DOI":"10.48550\/arXiv.2102.12149"},{"key":"9803_CR24","doi-asserted-by":"publisher","unstructured":"Voges, L. F., Jarren, L. C., & Seifert, S. (2023). Opening the random forest black box by the analysis of the mutual impact of features. https:\/\/doi.org\/10.48550\/arXiv.2304.02490","DOI":"10.48550\/arXiv.2304.02490"},{"key":"9803_CR25","doi-asserted-by":"publisher","unstructured":"Wang, J., C., Liu, C., Ouyang, Y., Qin, T., Lu, W., Chen, Y., Zeng, W., Philip, S., & Yu (2021). Generalizing to unseen domains: A survey on domain generalization. https:\/\/doi.org\/10.48550\/arXiv.2103.03097","DOI":"10.48550\/arXiv.2103.03097"},{"key":"9803_CR26","doi-asserted-by":"crossref","unstructured":"Yimam, M. S., Alemayehu, M. H., Ayele, A., & Biemann, C. (2020). Exploring Amharic sentiment analysis from social media texts: Building annotation tools and classification models. https:\/\/aclanthology.org\/2020.coling-main.91\/","DOI":"10.18653\/v1\/2020.coling-main.91"}],"container-title":["Language Resources and Evaluation"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10579-024-09803-2.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10579-024-09803-2\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10579-024-09803-2.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,7,25]],"date-time":"2025-07-25T12:26:36Z","timestamp":1753446396000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10579-024-09803-2"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,1,9]]},"references-count":26,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,9]]}},"alternative-id":["9803"],"URL":"https:\/\/doi.org\/10.1007\/s10579-024-09803-2","relation":{},"ISSN":["1574-020X","1574-0218"],"issn-type":[{"value":"1574-020X","type":"print"},{"value":"1574-0218","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,1,9]]},"assertion":[{"value":"16 December 2024","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"9 January 2025","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no competing interests.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interests"}},{"value":"Not commissioned; externally peer-reviewed.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Provenance and peer review"}}]}}