{"url":"/dataset/cad","name":"CAD","full_name":"Contextual Abuse Dataset","description_markdown":"Dataset of primarily English Reddit entries which addresses several limitations of prior work. It (1) contains six conceptually distinct primary categories as well as secondary categories, (2) has labels annotated in the context of the conversation thread, (3) contains rationales and (4) uses an expert-driven group-adjudication process for high quality annotations.","description_withheld":null,"homepage":"https://zenodo.org/record/4881008","introduced_date":"2021-06-01","introduced_date_note":null,"introduced_by":{"paper":"/paper/introducing-cad-the-contextual-abuse-dataset","title":"Introducing CAD: the Contextual Abuse Dataset","first_author":"Bertie Vidgen","url":null},"license":{"name":"Creative Commons Attribution 4.0 International Public License","url":"https://creativecommons.org/licenses/by/4.0/legalcode"},"modalities":[{"name":"Texts","url":"/datasets/modality/texts"}],"tasks":[{"name":"Toxic Comment Classification","url":"/task/toxic-comment-classification","datasets_with_task":"/datasets/task/toxic-comment-classification"},{"name":"Toxic Spans Detection","url":"/task/toxic-spans-detection","datasets_with_task":"/datasets/task/toxic-spans-detection"}],"languages":[{"name":"English","url":"/datasets/language/english"}],"variants":["CAD"],"data_loaders":[],"num_papers_in_archive":8,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/toxic-comment-classification-on-cad","task":"Toxic Comment Classification","dataset_variant":"CAD","rows":1,"metrics":["F1 on Toxic class"],"first_row_in_archive_order":{"model":"ContextRNN","paper":"/paper/revisiting-contextual-toxicity-detection-in","metrics":{"F1 on Toxic class":"52.5"},"code_links":[]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/revisiting-contextual-toxicity-detection-in","title":"Revisiting Contextual Toxicity Detection in Conversations","date":"2021-11-24","rows_on_this_dataset":1,"code_links":0,"syntology":null}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":0,"samples_harvested":0,"samples_ran":0,"samples_unverified":0,"pointer_only_for_licence":0,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}