{"url":"/task/answer-generation","name":"Answer Generation","slug":"answer-generation","description_markdown":null,"categories":[{"name":"Natural Language Processing","url":"/area/natural-language-processing"},{"name":"Reasoning","url":"/area/reasoning"}],"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28","slug_source":"archive_url"},"counts":{"papers_tagged":280,"papers_with_code":111,"benchmarks":2,"benchmark_tables_in_archive":2,"benchmark_tables_shown":2,"benchmark_tables_withheld_as_spam":0,"benchmark_definition":"a leaderboard table with at least one row; benchmark_tables_shown also counts the zero-row tables; benchmark_tables_in_archive adds the tables withheld as spam","datasets":4,"subtasks":0,"parent_tasks":1},"benchmarks":[{"leaderboard":"/sota/answer-generation-on-weibopolls","slug":"answer-generation-on-weibopolls","dataset":"WeiboPolls","dataset_url":"/dataset/weibopolls","rows_in_archive":3,"metrics":["ROUGE-1","ROUGE-L","BLEU-1","BLEU-3"],"first_row_in_archive_order":{"model":"UniPoll","paper_title":"UniPoll: A Unified Social Media Poll Generation Framework via Multi-Objective Optimization","paper_url":"/paper/unipoll-a-unified-social-media-poll","paper_date":"2023-06-12","arxiv_id":"2306.06851","code_links":[{"title":"X1AOX1A/UniPoll","url":"https://github.com/X1AOX1A/UniPoll"}],"syntology":null}},{"leaderboard":"/sota/answer-generation-on-cicero","slug":"answer-generation-on-cicero","dataset":"CICERO","dataset_url":"/dataset/cicero","rows_in_archive":2,"metrics":["ROUGE"],"first_row_in_archive_order":{"model":"T5-large pre-trained on GLUCOSE","paper_title":"CICERO: A Dataset for Contextualized Commonsense Inference in Dialogues","paper_url":"/paper/cicero-a-dataset-for-contextualized","paper_date":"2022-03-25","arxiv_id":"2203.13926","code_links":[{"title":"declare-lab/CICERO","url":"https://github.com/declare-lab/CICERO"}],"syntology":null}}],"datasets":[{"url":"/dataset/qampari","name":"QAMPARI","full_name":"","num_papers_in_archive":19},{"url":"/dataset/cicero","name":"CICERO","full_name":"Contextualized Commonsense Inference in Dialogues","num_papers_in_archive":12},{"url":"/dataset/weibopolls","name":"WeiboPolls","full_name":"","num_papers_in_archive":3},{"url":"/dataset/llmafia","name":"LLMafia","full_name":"","num_papers_in_archive":1}],"subtasks":[],"parent_tasks":[{"url":"/task/natural-language-inference","name":"Natural Language Inference"}],"papers":{"order":"repositories listed in the archive (desc), then date (desc); the archive holds no stars","population":"papers tagged with this task that list at least one repository in the archive","shown":30,"of":111,"tagged_in_all":280,"items":[{"url":"/paper/exploring-the-limits-of-transfer-learning","title":"Exploring the Limits of Transfer Learning with a Unified Text-to-Text Transformer","date":"2019-10-23","arxiv_id":"1910.10683","repositories_listed":57,"syntology":{"n":31,"n_ran":2,"n_unverified":29,"n_pointer_only":0}},{"url":"/paper/regnlp-in-action-facilitating-compliance","title":"RIRAG: Regulatory Information Retrieval and Answer Generation","date":"2024-09-09","arxiv_id":"2409.05677","repositories_listed":4,"syntology":null},{"url":"/paper/zusammenqa-data-augmentation-with-specialized","title":"ZusammenQA: Data Augmentation with Specialized Models for Cross-lingual Open-retrieval Question Answering System","date":"2022-05-30","arxiv_id":"2205.14981","repositories_listed":3,"syntology":null},{"url":"/paper/vogue-answer-verbalization-through-multi-task","title":"VOGUE: Answer Verbalization through Multi-Task Learning","date":"2021-06-24","arxiv_id":"2106.13316","repositories_listed":3,"syntology":null},{"url":"/paper/lrta-a-transparent-neural-symbolic-reasoning","title":"LRTA: A Transparent Neural-Symbolic Reasoning Framework with Modular Supervision for Visual Question Answering","date":"2020-11-21","arxiv_id":"2011.10731","repositories_listed":3,"syntology":null},{"url":"/paper/dataset-and-neural-recurrent-sequence","title":"Dataset and Neural Recurrent Sequence Labeling Model for Open-Domain Factoid Question Answering","date":"2016-07-21","arxiv_id":"1607.06275","repositories_listed":3,"syntology":null},{"url":"/paper/mm-instruct-generated-visual-instructions-for","title":"MM-Instruct: Generated Visual Instructions for Large Multimodal Model Alignment","date":"2024-06-28","arxiv_id":"2406.19736","repositories_listed":2,"syntology":null},{"url":"/paper/poisonedrag-knowledge-poisoning-attacks-to","title":"PoisonedRAG: Knowledge Corruption Attacks to Retrieval-Augmented Generation of Large Language Models","date":"2024-02-12","arxiv_id":"2402.07867","repositories_listed":2,"syntology":{"n":17,"n_ran":6,"n_unverified":11,"n_pointer_only":0}},{"url":"/paper/surgical-vqla-transformer-with-gated-vision","title":"Surgical-VQLA: Transformer with Gated Vision-Language Embedding for Visual Question Localized-Answering in Robotic Surgery","date":"2023-05-19","arxiv_id":"2305.11692","repositories_listed":2,"syntology":{"n":7,"n_ran":7,"n_unverified":0,"n_pointer_only":0}},{"url":"/paper/qampari-an-open-domain-question-answering","title":"QAMPARI: An Open-domain Question Answering Benchmark for Questions with Many Answers from Multiple Paragraphs","date":"2022-05-25","arxiv_id":"2205.12665","repositories_listed":2,"syntology":{"n":1,"n_ran":1,"n_unverified":0,"n_pointer_only":0}},{"url":"/paper/mumuqa-multimedia-multi-hop-news-question","title":"MuMuQA: Multimedia Multi-Hop News Question Answering via Cross-Media Knowledge Extraction and Grounding","date":"2021-12-20","arxiv_id":"2112.10728","repositories_listed":2,"syntology":null},{"url":"/paper/it-is-ai-s-turn-to-ask-human-a-question","title":"It is AI's Turn to Ask Humans a Question: Question-Answer Pair Generation for Children's Story Books","date":"2021-09-08","arxiv_id":"2109.03423","repositories_listed":2,"syntology":null},{"url":"/paper/end-to-end-training-of-multi-document-reader","title":"End-to-End Training of Multi-Document Reader and Retriever for Open-Domain Question Answering","date":"2021-06-09","arxiv_id":"2106.05346","repositories_listed":2,"syntology":{"n":4,"n_ran":4,"n_unverified":0,"n_pointer_only":4}},{"url":"/paper/quiz-style-question-generation-for-news","title":"Quiz-Style Question Generation for News Stories","date":"2021-02-18","arxiv_id":"2102.09094","repositories_listed":2,"syntology":null},{"url":"/paper/asking-questions-the-human-way-scalable","title":"Asking Questions the Human Way: Scalable Question-Answer Generation from Text Corpus","date":"2020-01-27","arxiv_id":"2002.00748","repositories_listed":2,"syntology":{"n":8,"n_ran":1,"n_unverified":7,"n_pointer_only":2}},{"url":"/paper/rmit-adm-s-at-the-sigir-2025-liverag","title":"RMIT-ADM+S at the SIGIR 2025 LiveRAG Challenge","date":"2025-06-17","arxiv_id":"2506.14516","repositories_listed":1,"syntology":null},{"url":"/paper/tablerag-a-retrieval-augmented-generation","title":"TableRAG: A Retrieval Augmented Generation Framework for Heterogeneous Document Reasoning","date":"2025-06-12","arxiv_id":"2506.10380","repositories_listed":1,"syntology":null},{"url":"/paper/ecorag-evidentiality-guided-compression-for","title":"ECoRAG: Evidentiality-guided Compression for Long Context RAG","date":"2025-06-05","arxiv_id":"2506.05167","repositories_listed":1,"syntology":null},{"url":"/paper/o-2-searcher-a-searching-based-agent-model","title":"O$^2$-Searcher: A Searching-based Agent Model for Open-Domain Open-Ended Question Answering","date":"2025-05-22","arxiv_id":"2505.16582","repositories_listed":1,"syntology":{"n":2,"n_ran":0,"n_unverified":2,"n_pointer_only":0}},{"url":"/paper/gui-g1-understanding-r1-zero-like-training","title":"GUI-G1: Understanding R1-Zero-Like Training for Visual Grounding in GUI Agents","date":"2025-05-21","arxiv_id":"2505.15810","repositories_listed":1,"syntology":{"n":10,"n_ran":1,"n_unverified":9,"n_pointer_only":0}},{"url":"/paper/the-atlas-of-in-context-learning-how","title":"The Atlas of In-Context Learning: How Attention Heads Shape In-Context Retrieval Augmentation","date":"2025-05-21","arxiv_id":"2505.15807","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_unverified":1,"n_pointer_only":1}},{"url":"/paper/process-vs-outcome-reward-which-is-better-for","title":"Process vs. Outcome Reward: Which is Better for Agentic RAG Reinforcement Learning","date":"2025-05-20","arxiv_id":"2505.14069","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_unverified":0,"n_pointer_only":1}},{"url":"/paper/kg-qagen-a-knowledge-graph-based-framework","title":"KG-QAGen: A Knowledge-Graph-Based Framework for Systematic Question Generation and Long-Context LLM Evaluation","date":"2025-05-18","arxiv_id":"2505.12495","repositories_listed":1,"syntology":null},{"url":"/paper/rag-vr-leveraging-retrieval-augmented","title":"RAG-VR: Leveraging Retrieval-Augmented Generation for 3D Question Answering in VR Environments","date":"2025-04-11","arxiv_id":"2504.08256","repositories_listed":1,"syntology":null},{"url":"/paper/localized-definitions-and-distributed","title":"Localized Definitions and Distributed Reasoning: A Proof-of-Concept Mechanistic Interpretability Study via Activation Patching","date":"2025-04-03","arxiv_id":"2504.02976","repositories_listed":1,"syntology":null},{"url":"/paper/metaladder-ascending-mathematical-solution","title":"MetaLadder: Ascending Mathematical Solution Quality via Analogical-Problem Reasoning Transfer","date":"2025-03-19","arxiv_id":"2503.14891","repositories_listed":1,"syntology":null},{"url":"/paper/conversational-gold-evaluating-personalized","title":"Conversational Gold: Evaluating Personalized Conversational Search System using Gold Nuggets","date":"2025-03-12","arxiv_id":"2503.09902","repositories_listed":1,"syntology":null},{"url":"/paper/zero-shot-complex-question-answering-on-long","title":"Zero-Shot Complex Question-Answering on Long Scientific Documents","date":"2025-03-04","arxiv_id":"2503.02695","repositories_listed":1,"syntology":null},{"url":"/paper/egonormia-benchmarking-physical-social-norm","title":"EgoNormia: Benchmarking Physical Social Norm Understanding","date":"2025-02-27","arxiv_id":"2502.20490","repositories_listed":1,"syntology":{"n":3,"n_ran":0,"n_unverified":3,"n_pointer_only":0}},{"url":"/paper/a-hybrid-approach-to-information-retrieval","title":"A Hybrid Approach to Information Retrieval and Answer Generation for Regulatory Texts","date":"2025-02-24","arxiv_id":"2502.16767","repositories_listed":1,"syntology":null}],"syntology_records":11,"syntology_note":"a paper without a record is not a recorded non-run: it may lack an arXiv id or simply be absent from the graph layer"},"description_links":{"kept":0,"unwrapped_to_text":0,"bare_urls_linked":0,"relative_images_dropped":0,"rule":"internal links are kept only when the target slug exists in the catalog"},"syntology":{"read_at":"2026-09-24T18:15:14+00:00","claim":"Per-sample execution status on synthesized fixtures ('ran N of M samples'); not a correctness claim and not a ranking signal.","status_vocabulary":{"ran_honours":"ran, honoured the contract we drafted","ran_violates":"ran, violated the contract we drafted","ran_draft_wrong":"ran; our contract draft was wrong, not the code","ran_fixture":"ran; our fixture could not drive it","ran":"ran on a synthesized input","unverified":"unverified (harvested, no recorded run)"}},"not_shown":{"libraries":"the archive has no per-task library table","trend_sparklines":"the Trend column of the benchmarks table was a rendered image; it is not in the archive","social_and_latest_sorts":"stars and social signals are not in the archive"}}