{"url":"/dataset/sarc","name":"SARC","full_name":null,"description_markdown":"This dataset was designed for contextual investigations, with related works making considerable usage of said context. The dataset was constructed by scraping Reddit comments; with sarcastic entries being self-annotated by authors through the use of the \\s token, which indicates sarcastic intent on the website. Posts on Reddit are often in response to another comment; SARC incorporates this information through the addition of the parent comment and further child comments surrounding a post. \r\n\r\nSource: [DEEP AND DENSE SARCASM DETECTION](https://arxiv.org/pdf/1911.07474v2.pdf)","description_withheld":null,"homepage":"https://nlp.cs.princeton.edu/SARC/","introduced_date":"2017-04-19","introduced_date_note":null,"introduced_by":{"paper":"/paper/a-large-self-annotated-corpus-for-sarcasm","title":"A Large Self-Annotated Corpus for Sarcasm","first_author":"Mikhail Khodak","url":null},"license":{"name":"Unknown","url":null},"modalities":[],"tasks":[{"name":"Sarcasm Detection","url":"/task/sarcasm-detection","datasets_with_task":"/datasets/task/sarcasm-detection"}],"languages":[],"variants":["SARC (all-bal)","SARC (pol-bal)","SARC (pol-unbal)","SARC"],"data_loaders":[{"repo":"https://github.com/harsh252/Sarcasm_detection","url":"https://github.com/harsh252/Sarcasm_detection","frameworks":[]}],"num_papers_in_archive":35,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/sarcasm-detection-on-sarc-all-bal","task":"Sarcasm Detection","dataset_variant":"SARC (all-bal)","rows":2,"metrics":["Accuracy"],"first_row_in_archive_order":{"model":"CASCADE","paper":"/paper/cascade-contextual-sarcasm-detection-in","metrics":{"Accuracy":"77"},"code_links":[{"title":"SenticNet/CASCADE--ContextuAl-SarCAsm-DEtector","url":"https://github.com/SenticNet/CASCADE--ContextuAl-SarCAsm-DEtector"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/sarcasm-detection-on-sarc-pol-bal","task":"Sarcasm Detection","dataset_variant":"SARC (pol-bal)","rows":2,"metrics":["Accuracy"],"first_row_in_archive_order":{"model":"Bag-of-Bigrams","paper":"/paper/a-large-self-annotated-corpus-for-sarcasm","metrics":{"Accuracy":"76.5"},"code_links":[{"title":"NLPrinceton/SARC","url":"https://github.com/NLPrinceton/SARC"},{"title":"chrisolen1/sarcasm-detection","url":"https://github.com/chrisolen1/sarcasm-detection"},{"title":"NauqGnesh/RedditSarcasm","url":"https://github.com/NauqGnesh/RedditSarcasm"},{"title":"Kaguura/SarcasmDetection","url":"https://github.com/Kaguura/SarcasmDetection"},{"title":"karlwbaker/Springboard_capstone","url":"https://github.com/karlwbaker/Springboard_capstone"},{"title":"sachinsharma3191/Sarcasm-Detection","url":"https://github.com/sachinsharma3191/Sarcasm-Detection"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/sarcasm-detection-on-sarc-pol-unbal","task":"Sarcasm Detection","dataset_variant":"SARC (pol-unbal)","rows":1,"metrics":["Avg F1"],"first_row_in_archive_order":{"model":"Bag-of-Words","paper":"/paper/a-large-self-annotated-corpus-for-sarcasm","metrics":{"Avg F1":"27.0"},"code_links":[{"title":"NLPrinceton/SARC","url":"https://github.com/NLPrinceton/SARC"},{"title":"chrisolen1/sarcasm-detection","url":"https://github.com/chrisolen1/sarcasm-detection"},{"title":"NauqGnesh/RedditSarcasm","url":"https://github.com/NauqGnesh/RedditSarcasm"},{"title":"Kaguura/SarcasmDetection","url":"https://github.com/Kaguura/SarcasmDetection"},{"title":"karlwbaker/Springboard_capstone","url":"https://github.com/karlwbaker/Springboard_capstone"},{"title":"sachinsharma3191/Sarcasm-Detection","url":"https://github.com/sachinsharma3191/Sarcasm-Detection"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/cascade-contextual-sarcasm-detection-in","title":"CASCADE: Contextual Sarcasm Detection in Online Discussion Forums","date":"2018-05-16","rows_on_this_dataset":2,"code_links":1,"syntology":null},{"paper":"/paper/a-large-self-annotated-corpus-for-sarcasm","title":"A Large Self-Annotated Corpus for Sarcasm","date":"2017-04-19","rows_on_this_dataset":3,"code_links":6,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":1,"samples_ran":1,"samples_unverified":0,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":1,"samples_harvested":1,"samples_ran":1,"samples_unverified":0,"pointer_only_for_licence":0,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}