@inproceedings{reshma-etal-2023-mitigating,
title = "Mitigating Abusive Comment Detection in {T}amil Text: A Data Augmentation Approach with Transformer Model",
author = "Reshma, Sheik and
Raghavan, Balanathan and
Jaya Nirmala, S.",
editor = "Jyoti, D. Pawar and
Sobha, Lalitha Devi",
booktitle = "Proceedings of the 20th International Conference on Natural Language Processing (ICON)",
month = dec,
year = "2023",
address = "Goa University, Goa, India",
publisher = "NLP Association of India (NLPAI)",
url = "https://aclanthology.org/2023.icon-1.39",
pages = "460--465",
abstract = "With the increasing number of users on social media platforms, the detection and categorization of abusive comments have become crucial, necessitating effective strategies to mitigate their impact on online discussions. However, the intricate and diverse nature of lowresource Indic languages presents a challenge in developing reliable detection methodologies. This research focuses on the task of classifying YouTube comments written in Tamil language into various categories. To achieve this, our research conducted experiments utilizing various multi-lingual transformer-based models along with data augmentation approaches involving back translation approaches and other pre-processing techniques. Our work provides valuable insights into the effectiveness of various preprocessing methods for this classification task. Our experiments showed that the Multilingual Representations for Indian Languages (MURIL) transformer model, coupled with round-trip translation and lexical replacement, yielded the most promising results, showcasing a significant improvement of over 15 units in macro F1-score compared to existing baselines. This contribution adds to the ongoing research to mitigate the adverse impact of abusive content on online platforms, emphasizing the utilization of diverse preprocessing strategies and state-of-the-art language models.",
}
-
Notifications
You must be signed in to change notification settings - Fork 0
π [ICON 2023] "Mitigating Abusive Comment Detection in Tamil Text: A Data Augmentation Approach with Transformer Model", Reshma Sheik and Raghavan Balanathan and Dr. S. Jaya Nirmala
License
ICON-ML/ACDT
Folders and files
Name | Name | Last commit message | Last commit date | |
---|---|---|---|---|
Β | Β | |||
Β | Β | |||
Β | Β | |||
Β | Β | |||
Repository files navigation
About
π [ICON 2023] "Mitigating Abusive Comment Detection in Tamil Text: A Data Augmentation Approach with Transformer Model", Reshma Sheik and Raghavan Balanathan and Dr. S. Jaya Nirmala
Resources
License
Stars
Watchers
Forks
Releases
No releases published
Packages 0
No packages published