{"url":"/dataset/bengali-cyberbullying-detection-comments","name":"Bengali Cyberbullying Detection Comments Dataset","full_name":null,"description_markdown":"This dataset contains 44,001 Bengali comments, curated to detect cyberbullying using Natural Language Processing (NLP) techniques. Each comment is labeled by experts, categorizing different forms of harassment and offensive behavior. The dataset enables the identification of inappropriate content, ranging from mild to severe harassment, facilitating precise classification and analysis. This resource is designed for researchers and developers working on cyberbullying detection, sentiment analysis, and content moderation in Bengali text.\r\n\r\nIf you want to work further on this dataset, cite this paper - https://arxiv.org/abs/2106.04506","description_withheld":null,"homepage":"https://www.kaggle.com/datasets/cypher1337/dataset-for-cyberbully-detection-bengali-comments","introduced_date":null,"introduced_date_note":null,"introduced_by":null,"license":{"name":"Database: Open Database, Contents: © Original Authors","url":null},"modalities":[],"tasks":[],"languages":[{"name":"Bengali","url":"/datasets/language/bengali"}],"variants":["Bengali Cyberbullying Detection Comments Dataset"],"data_loaders":[],"num_papers_in_archive":0,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[],"papers_with_a_benchmark_row":[],"syntology_totals":{"read_at":"2026-09-25T09:33:49+00:00","papers_with_samples":0,"samples_harvested":0,"samples_ran":0,"samples_unverified":0,"pointer_only_for_licence":0,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}