{"url":"/dataset/instructional-dt-instr-dt","name":"Instructional-DT (Instr-DT)","full_name":"Instructional Discourse Treebank","description_markdown":"This discourse treebank includes annotated instructional texts originally assembled at the Information Technology Research Institute, University of Brighton. This dataset contains 176 documents with an average of 32.6 EDUs for a total of 5744 EDUs and 53,250 words.","description_withheld":null,"homepage":"","introduced_date":"2009-05-31","introduced_date_note":null,"introduced_by":{"paper":"/paper/an-effective-discourse-parser-that-uses-rich","title":"An effective Discourse Parser that uses Rich Linguistic Information","first_author":"Rajen Subba","url":null},"license":null,"modalities":[{"name":"Texts","url":"/datasets/modality/texts"}],"tasks":[{"name":"Discourse Parsing","url":"/task/discourse-parsing","datasets_with_task":"/datasets/task/discourse-parsing"}],"languages":[{"name":"English","url":"/datasets/language/english"}],"variants":["Instructional-DT (Instr-DT)"],"data_loaders":[],"num_papers_in_archive":5,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/discourse-parsing-on-instructional-dt-instr","task":"Discourse Parsing","dataset_variant":"Instructional-DT (Instr-DT)","rows":12,"metrics":["Standard Parseval (Nuclearity)","Standard Parseval (Span)","Standard Parseval (Full)","Standard Parseval (Relation)"],"first_row_in_archive_order":{"model":"Bottom-up (DeBERTa)","paper":"/paper/a-simple-and-strong-baseline-for-end-to-end","metrics":{"Standard Parseval (Full)":"44.4","Standard Parseval (Nuclearity)":"60.0","Standard Parseval (Relation)":"51.4","Standard Parseval (Span)":"77.8"},"code_links":[{"title":"nttcslab-nlp/rstparser_emnlp22","url":"https://github.com/nttcslab-nlp/rstparser_emnlp22"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/a-simple-and-strong-baseline-for-end-to-end","title":"A Simple and Strong Baseline for End-to-End Neural RST-style Discourse Parsing","date":"2022-10-15","rows_on_this_dataset":10,"code_links":1,"syntology":null},{"paper":"/paper/unleashing-the-power-of-neural-discourse","title":"Unleashing the Power of Neural Discourse Parsers -- A Context and Structure Aware Approach Using Large Scale Pretraining","date":"2020-11-06","rows_on_this_dataset":2,"code_links":0,"syntology":null}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":0,"samples_harvested":0,"samples_ran":0,"samples_unverified":0,"pointer_only_for_licence":0,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}