{"url":"/dataset/arxiv","name":"Arxiv HEP-TH citation graph","full_name":null,"description_markdown":"**Arxiv HEP-TH (high energy physics theory) citation graph** is from the e-print **arXiv** and covers all the citations within a dataset of 27,770 papers with 352,807 edges. If a paper i cites paper j, the graph contains a directed edge from i to j. If a paper cites, or is cited by, a paper outside the dataset, the graph does not contain any information about this.\r\nThe data covers papers in the period from January 1993 to April 2003 (124 months).\r\n\r\nSource: [https://snap.stanford.edu/data/cit-HepTh.html](https://snap.stanford.edu/data/cit-HepTh.html)","description_withheld":null,"homepage":"https://snap.stanford.edu/data/cit-HepTh.html","introduced_date":null,"introduced_date_note":null,"introduced_by":null,"license":null,"modalities":[{"name":"Graphs","url":"/datasets/modality/graphs"}],"tasks":[{"name":"Text Classification","url":"/task/text-classification","datasets_with_task":"/datasets/task/text-classification"},{"name":"Text Summarization","url":"/task/text-summarization","datasets_with_task":"/datasets/task/text-summarization"},{"name":"Language Modelling","url":"/task/language-modelling","datasets_with_task":"/datasets/task/language-modelling"},{"name":"Topic Models","url":"/task/topic-models","datasets_with_task":"/datasets/task/topic-models"},{"name":"Document Summarization","url":"/task/document-summarization","datasets_with_task":"/datasets/task/document-summarization"},{"name":"Extended Summarization","url":"/task/extended-summarization","datasets_with_task":"/datasets/task/extended-summarization"},{"name":"Clique Prediction","url":"/task/clique-prediction","datasets_with_task":"/datasets/task/clique-prediction"}],"languages":[],"variants":["arXiv-AstroPh 3-clique","arXiv-GrQc 4-clique","Arxiv HEP-TH citation graph","arXiv-Long Val","arXiv-Long Test"],"data_loaders":[],"num_papers_in_archive":35,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/text-summarization-on-arxiv","task":"Text Summarization","dataset_variant":"Arxiv HEP-TH citation graph","rows":28,"metrics":["ROUGE-1","ROUGE-2","ROUGE-L"],"first_row_in_archive_order":{"model":"Top Down Transformer (AdaPool) (464M)","paper":"/paper/long-document-summarization-with-top-down-and-1","metrics":{"ROUGE-1":"50.95","ROUGE-2":"21.93","ROUGE-L":"45.61"},"code_links":[{"title":"lllyasviel/framepack","url":"https://github.com/lllyasviel/framepack"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/topic-models-on-arxiv","task":"Topic Models","dataset_variant":"Arxiv HEP-TH citation graph","rows":2,"metrics":["MACC","Topic Coherence@50","Topic coherence@5"],"first_row_in_archive_order":{"model":"JoSH","paper":"/paper/hierarchical-topic-mining-via-joint-spherical","metrics":{"MACC":"83.24","Topic coherence@5":"0.0074"},"code_links":[{"title":"yumeng5/JoSH","url":"https://github.com/yumeng5/JoSH"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/document-summarization-on-arxiv","task":"Document Summarization","dataset_variant":"Arxiv HEP-TH citation graph","rows":1,"metrics":["ROUGE-1"],"first_row_in_archive_order":{"model":"DeepPyramidion","paper":"/paper/sparsifying-transformer-models-with-trainable","metrics":{"ROUGE-1":"47.15"},"code_links":[]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/language-modelling-on-arxiv","task":"Language Modelling","dataset_variant":"Arxiv HEP-TH citation graph","rows":1,"metrics":["BPB"],"first_row_in_archive_order":{"model":"Gopher","paper":"/paper/scaling-language-models-methods-analysis-1","metrics":{"BPB":"0.662"},"code_links":[{"title":"allenai/dolma","url":"https://github.com/allenai/dolma"},{"title":"rvlopes/gloria","url":"https://github.com/rvlopes/gloria"},{"title":"bramiozo/PubScience","url":"https://github.com/bramiozo/PubScience"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/text-classification-on-arxiv","task":"Text Classification","dataset_variant":"Arxiv HEP-TH citation graph","rows":1,"metrics":["Accuracy"],"first_row_in_archive_order":{"model":"BigBird","paper":"/paper/big-bird-transformers-for-longer-sequences","metrics":{"Accuracy":"92.31"},"code_links":[{"title":"huggingface/transformers","url":"https://github.com/huggingface/transformers"},{"title":"tensorflow/models","url":"https://github.com/tensorflow/models/tree/master/official/nlp/projects/bigbird"},{"title":"PaddlePaddle/PaddleNLP","url":"https://github.com/PaddlePaddle/PaddleNLP/tree/develop/paddlenlp/transformers/bigbird"},{"title":"facebookresearch/xformers","url":"https://github.com/facebookresearch/xformers"},{"title":"google-research/bigbird","url":"https://github.com/google-research/bigbird"},{"title":"monologg/kobigbird","url":"https://github.com/monologg/kobigbird"},{"title":"mim-solutions/bert_for_longer_texts","url":"https://github.com/mim-solutions/bert_for_longer_texts"},{"title":"mim-solutions/roberta_for_longer_texts","url":"https://github.com/mim-solutions/roberta_for_longer_texts"},{"title":"sajjjadayobi/ParsBigBird","url":"https://github.com/sajjjadayobi/ParsBigBird"},{"title":"thefonseca/factorsum","url":"https://github.com/thefonseca/factorsum"},{"title":"2024-MindSpore-1/Code2","url":"https://github.com/2024-MindSpore-1/Code2/tree/main/model-1/big_bird"},{"title":"sergeykramp/mthesis-bigbird-embeddings","url":"https://github.com/sergeykramp/mthesis-bigbird-embeddings"},{"title":"pwc-1/Paper-8","url":"https://github.com/pwc-1/Paper-8/tree/main/big_bird"},{"title":"pwc-1/Paper-8","url":"https://github.com/pwc-1/Paper-8/tree/main/bigbird_pegasus"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/segmented-recurrent-transformer-an-efficient","title":"Segmented Recurrent Transformer: An Efficient Sequence-to-Sequence Model","date":"2023-05-24","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/document-summarization-with-text-segmentation","title":"Document Summarization with Text Segmentation","date":"2023-01-20","rows_on_this_dataset":2,"code_links":0,"syntology":null},{"paper":"/paper/toward-unifying-text-segmentation-and-long","title":"Toward Unifying Text Segmentation and Long Document Summarization","date":"2022-10-28","rows_on_this_dataset":2,"code_links":1,"syntology":null},{"paper":"/paper/adapting-pretrained-text-to-text-models-for","title":"Adapting Pretrained Text-to-Text Models for Long Text Sequences","date":"2022-09-21","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/investigating-efficiently-extending","title":"Investigating Efficiently Extending Transformers for Long Input Summarization","date":"2022-08-08","rows_on_this_dataset":1,"code_links":2,"syntology":null},{"paper":"/paper/factorizing-content-and-budget-decisions-in","title":"Factorizing Content and Budget Decisions in Abstractive Summarization of Long Documents","date":"2022-05-25","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/gencomparesum-a-hybrid-unsupervised","title":"GenCompareSum: a hybrid unsupervised summarization method using salience","date":"2022-05-01","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/histruct-improving-extractive-text-1","title":"HiStruct+: Improving Extractive Text Summarization with Hierarchical Structure Information","date":"2022-03-17","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/long-document-summarization-with-top-down-and-1","title":"Long Document Summarization with Top-down and Bottom-up Inference","date":"2022-03-15","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/longt5-efficient-text-to-text-transformer-for","title":"LongT5: Efficient Text-To-Text Transformer for Long Sequences","date":"2021-12-15","rows_on_this_dataset":1,"code_links":4,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":1,"samples_ran":1,"samples_unverified":0,"pointer_only_for_licence":1,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/scaling-language-models-methods-analysis-1","title":"Scaling Language Models: Methods, Analysis & Insights from Training Gopher","date":"2021-12-08","rows_on_this_dataset":1,"code_links":3,"syntology":null},{"paper":"/paper/sparsifying-transformer-models-with-trainable","title":"Sparsifying Transformer Models with Trainable Representation Pooling","date":"2021-11-16","rows_on_this_dataset":3,"code_links":0,"syntology":null},{"paper":"/paper/memsum-extractive-summarization-of-long","title":"MemSum: Extractive Summarization of Long Documents Using Multi-Step Episodic Markov Decision Processes","date":"2021-07-19","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/hierarchical-learning-for-generation-with","title":"Hierarchical Learning for Generation with Long Source Sequences","date":"2021-04-15","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/systematically-exploring-redundancy-reduction","title":"Systematically Exploring Redundancy Reduction in Summarizing Long Documents","date":"2020-11-30","rows_on_this_dataset":2,"code_links":1,"syntology":null},{"paper":"/paper/big-bird-transformers-for-longer-sequences","title":"Big Bird: Transformers for Longer Sequences","date":"2020-07-28","rows_on_this_dataset":1,"code_links":14,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":15,"samples_ran":10,"samples_unverified":5,"pointer_only_for_licence":11,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/hierarchical-topic-mining-via-joint-spherical","title":"Hierarchical Topic Mining via Joint Spherical Tree and Text Embedding","date":"2020-07-18","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/a-divide-and-conquer-approach-to-the","title":"A Divide-and-Conquer Approach to the Summarization of Long Documents","date":"2020-04-13","rows_on_this_dataset":3,"code_links":1,"syntology":null},{"paper":"/paper/pegasus-pre-training-with-extracted-gap","title":"PEGASUS: Pre-training with Extracted Gap-sentences for Abstractive Summarization","date":"2019-12-18","rows_on_this_dataset":1,"code_links":19,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":19,"samples_ran":1,"samples_unverified":18,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/extractive-summarization-of-long-documents-by","title":"Extractive Summarization of Long Documents by Combining Global and Local Context","date":"2019-09-17","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/on-extractive-and-abstractive-neural-document","title":"On Extractive and Abstractive Neural Document Summarization with Transformer Language Models","date":"2019-09-07","rows_on_this_dataset":3,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":6,"samples_ran":0,"samples_unverified":6,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/topiceq-a-joint-topic-and-mathematical","title":"TopicEq: A Joint Topic and Mathematical Equation Model for Scientific Texts","date":"2019-02-16","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/a-discourse-aware-attention-model-for","title":"A Discourse-Aware Attention Model for Abstractive Summarization of Long Documents","date":"2018-04-16","rows_on_this_dataset":1,"code_links":3,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":14,"samples_ran":1,"samples_unverified":13,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/get-to-the-point-summarization-with-pointer","title":"Get To The Point: Summarization with Pointer-Generator Networks","date":"2017-04-14","rows_on_this_dataset":1,"code_links":39,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":64,"samples_ran":30,"samples_unverified":34,"pointer_only_for_licence":44,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":6,"samples_harvested":119,"samples_ran":43,"samples_unverified":76,"pointer_only_for_licence":56,"papers_with_no_sample_that_ran":1,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}