{"url":"/dataset/bcopa-ce","name":"BCOPA-CE","full_name":"A Balanced COPA Test Set with cause-effect as alternatives","description_markdown":"We provide the BCOPA-CE test set, which has balanced token distribution in the correct and wrong alternatives and increases the difficulty of being aware of cause and effect.\r\n### construction \r\n1. for each premise of the 500 samples in COPA-test set, we generate one event manually which is a plausible answer to the opposite question type of the original sample.\r\n2. obtain 500 triplets of <*premise*, *cause*, *effect*>\r\n3. construct 1000 samples by giving two different questions (**cause** or **effect**) to each triplet.","description_withheld":null,"homepage":"https://github.com/badbadcode/weakCOPA/blob/master/data/BCOPA-CE.xml","introduced_date":"2021-07-05","introduced_date_note":null,"introduced_by":{"paper":"/paper/doing-good-or-doing-right-exploring-the","title":"Doing Good or Doing Right? Exploring the Weakness of Commonsense Causal Reasoning Models","first_author":"Mingyue Han","url":null},"license":null,"modalities":[{"name":"Texts","url":"/datasets/modality/texts"}],"tasks":[{"name":"Common Sense Reasoning","url":"/task/common-sense-reasoning","datasets_with_task":"/datasets/task/common-sense-reasoning"},{"name":"Causal Identification","url":"/task/causal-identification","datasets_with_task":"/datasets/task/causal-identification"},{"name":"Causal Discovery","url":"/task/causal-discovery","datasets_with_task":"/datasets/task/causal-discovery"}],"languages":[{"name":"English","url":"/datasets/language/english"}],"variants":["BCOPA-CE"],"data_loaders":[],"num_papers_in_archive":3,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[],"papers_with_a_benchmark_row":[],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":0,"samples_harvested":0,"samples_ran":0,"samples_unverified":0,"pointer_only_for_licence":0,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}