{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/cosine-annealing/papers/31","list_of":"/method/cosine-annealing","method":"Cosine Annealing","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":31,"pages_in_order":40,"rows_per_page":100,"rows":[3001,3100],"of":3965,"counts":{"archive_papers_tagged":3965,"with_a_code_link":1734,"where_syntology_ran_a_sample":627,"not_listed_spam_title":0,"listed":3965,"listed_where_code_ran":627,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":513,"every_run_a_failure_of_syntologys_instrument":114,"listed_with_a_run_with_no_instrument_failure":513,"listed_every_run_a_failure_of_syntologys_instrument":114,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/cosine-annealing","prev":"/method/cosine-annealing/papers/30","next":"/method/cosine-annealing/papers/32","papers":[{"paper":"/paper/llm-planner-few-shot-grounded-planning-for","slug":"llm-planner-few-shot-grounded-planning-for","title":"LLM-Planner: Few-Shot Grounded Planning for Embodied Agents with Large Language Models","date":"2022-12-08","arxiv_id":"2212.04088","n_code_links":1,"syntology":null},{"paper":"/paper/np4g-network-programming-for-generalization","slug":"np4g-network-programming-for-generalization","title":"NP4G : Network Programming for Generalization","date":"2022-12-08","arxiv_id":"2212.11118","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-role-of-ai-in-drug-discovery-challenges","title":"The Role of AI in Drug Discovery: Challenges, Opportunities, and Strategies","date":"2022-12-08","arxiv_id":"2212.08104","n_code_links":0,"syntology":null},{"paper":"/paper/deepspeed-data-efficiency-improving-deep","slug":"deepspeed-data-efficiency-improving-deep","title":"DeepSpeed Data Efficiency: Improving Deep Learning Model Quality and Training Efficiency via Efficient Data Sampling and Routing","date":"2022-12-07","arxiv_id":"2212.03597","n_code_links":1,"syntology":null},{"paper":"/paper/adaptive-testing-of-computer-vision-models","slug":"adaptive-testing-of-computer-vision-models","title":"Adaptive Testing of Computer Vision Models","date":"2022-12-06","arxiv_id":"2212.02774","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":4,"n_instrument":0,"unverified":2,"pointer_only":6,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["i-gao/adavision"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/counterfactual-reasoning-do-language-models","slug":"counterfactual-reasoning-do-language-models","title":"Counterfactual reasoning: Do language models need world knowledge for causal understanding?","date":"2022-12-06","arxiv_id":"2212.03278","n_code_links":1,"syntology":null},{"paper":null,"slug":"mobileptx-sparse-coding-for-pneumothorax","title":"MobilePTX: Sparse Coding for Pneumothorax Detection Given Limited Training Examples","date":"2022-12-06","arxiv_id":"2212.03282","n_code_links":0,"syntology":null},{"paper":null,"slug":"modern-french-poetry-generation-with-roberta","title":"Modern French Poetry Generation with RoBERTa and GPT-2","date":"2022-12-06","arxiv_id":"2212.02911","n_code_links":0,"syntology":null},{"paper":null,"slug":"audio-driven-co-speech-gesture-video","title":"Audio-Driven Co-Speech Gesture Video Generation","date":"2022-12-05","arxiv_id":"2212.02350","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-generation-of-factual-news","title":"Automatic Generation of Factual News Headlines in Finnish","date":"2022-12-05","arxiv_id":"2212.02170","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-the-limits-of-differentially","title":"Exploring the Limits of Differentially Private Deep Learning with Group-wise Clipping","date":"2022-12-03","arxiv_id":"2212.01539","n_code_links":0,"syntology":null},{"paper":"/paper/sumren-summarizing-reported-speech-about","slug":"sumren-summarizing-reported-speech-about","title":"SumREN: Summarizing Reported Speech about Events in News","date":"2022-12-02","arxiv_id":"2212.01146","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-survey-on-gpt-3","title":"a survey on GPT-3","date":"2022-12-01","arxiv_id":"2212.00857","n_code_links":0,"syntology":null},{"paper":"/paper/distilling-multi-step-reasoning-capabilities","slug":"distilling-multi-step-reasoning-capabilities","title":"Distilling Reasoning Capabilities into Smaller Language Models","date":"2022-12-01","arxiv_id":"2212.00193","n_code_links":1,"syntology":null},{"paper":null,"slug":"quadapter-adapter-for-gpt-2-quantization","title":"Quadapter: Adapter for GPT-2 Quantization","date":"2022-11-30","arxiv_id":"2211.16912","n_code_links":0,"syntology":null},{"paper":null,"slug":"outfit-generation-and-recommendation-an","title":"Outfit Generation and Recommendation -- An Experimental Study","date":"2022-11-29","arxiv_id":"2211.16353","n_code_links":0,"syntology":null},{"paper":"/paper/zero-shot-opinion-summarization-with-gpt-3","slug":"zero-shot-opinion-summarization-with-gpt-3","title":"Prompted Opinion Summarization with GPT-3.5","date":"2022-11-29","arxiv_id":"2211.15914","n_code_links":1,"syntology":null},{"paper":"/paper/gpt-neo-for-commonsense-reasoning-a","slug":"gpt-neo-for-commonsense-reasoning-a","title":"GPT-Neo for commonsense reasoning -- a theoretical and practical lens","date":"2022-11-28","arxiv_id":"2211.15593","n_code_links":1,"syntology":null},{"paper":"/paper/scientific-and-creative-analogies-in","slug":"scientific-and-creative-analogies-in","title":"Scientific and Creative Analogies in Pretrained Language Models","date":"2022-11-28","arxiv_id":"2211.15268","n_code_links":2,"syntology":null},{"paper":null,"slug":"understanding-bloom-an-empirical-study-on","title":"Understanding BLOOM: An empirical study on diverse NLP tasks","date":"2022-11-27","arxiv_id":"2211.14865","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-3-driven-pedagogical-agents-for-training","title":"GPT-3-driven pedagogical agents for training children's curious question-asking skills","date":"2022-11-25","arxiv_id":"2211.14228","n_code_links":0,"syntology":null},{"paper":"/paper/prompttts-controllable-text-to-speech-with","slug":"prompttts-controllable-text-to-speech-with","title":"PromptTTS: Controllable Text-to-Speech with Text Descriptions","date":"2022-11-22","arxiv_id":"2211.12171","n_code_links":1,"syntology":null},{"paper":"/paper/exploring-the-efficacy-of-pre-trained","slug":"exploring-the-efficacy-of-pre-trained","title":"Exploring the Efficacy of Pre-trained Checkpoints in Text-to-Music Generation Task","date":"2022-11-21","arxiv_id":"2211.11216","n_code_links":2,"syntology":null},{"paper":"/paper/language-in-a-bottle-language-model-guided","slug":"language-in-a-bottle-language-model-guided","title":"Language in a Bottle: Language Model Guided Concept Bottlenecks for Interpretable Image Classification","date":"2022-11-21","arxiv_id":"2211.11158","n_code_links":2,"syntology":null},{"paper":"/paper/pointclip-v2-adapting-clip-for-powerful-3d","slug":"pointclip-v2-adapting-clip-for-powerful-3d","title":"PointCLIP V2: Prompting CLIP and GPT for Powerful 3D Open-world Learning","date":"2022-11-21","arxiv_id":"2211.11682","n_code_links":2,"syntology":{"ran":8,"of":12,"n_ran_checked":4,"n_instrument":4,"unverified":4,"pointer_only":4,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 1 violated, 3 with no contract checked; 4 where Syntology's instrument failed) · 4 unverified","official":{"repos":["yangyangyang127/pointclip_v2"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"conceptor-aided-debiasing-of-contextualized","title":"Conceptor-Aided Debiasing of Large Language Models","date":"2022-11-20","arxiv_id":"2211.11087","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-knowledge-enhanced-multimodal","title":"A survey on knowledge-enhanced multimodal learning","date":"2022-11-19","arxiv_id":"2211.12328","n_code_links":0,"syntology":null},{"paper":"/paper/ignore-previous-prompt-attack-techniques-for","slug":"ignore-previous-prompt-attack-techniques-for","title":"Ignore Previous Prompt: Attack Techniques For Language Models","date":"2022-11-17","arxiv_id":"2211.09527","n_code_links":1,"syntology":{"ran":2,"of":6,"n_ran_checked":2,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["agencyenterprise/promptinject"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/random-ltd-random-and-layerwise-token","slug":"random-ltd-random-and-layerwise-token","title":"Random-LTD: Random and Layerwise Token Dropping Brings Efficient Training for Large-scale Transformers","date":"2022-11-17","arxiv_id":"2211.11586","n_code_links":1,"syntology":null},{"paper":"/paper/unisumm-unified-few-shot-summarization-with","slug":"unisumm-unified-few-shot-summarization-with","title":"UniSumm and SummZoo: Unified Model and Diverse Benchmark for Few-Shot Summarization","date":"2022-11-17","arxiv_id":"2211.09783","n_code_links":1,"syntology":null},{"paper":"/paper/galactica-a-large-language-model-for-science-1","slug":"galactica-a-large-language-model-for-science-1","title":"Galactica: A Large Language Model for Science","date":"2022-11-16","arxiv_id":"2211.09085","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["paperswithcode/galai"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"tsmind-alibaba-and-soochow-university-s","title":"TSMind: Alibaba and Soochow University's Submission to the WMT22 Translation Suggestion Task","date":"2022-11-16","arxiv_id":"2211.08987","n_code_links":0,"syntology":null},{"paper":"/paper/glue-x-evaluating-natural-language","slug":"glue-x-evaluating-natural-language","title":"GLUE-X: Evaluating Natural Language Understanding Models from an Out-of-distribution Generalization Perspective","date":"2022-11-15","arxiv_id":"2211.08073","n_code_links":1,"syntology":null},{"paper":"/paper/promptcap-prompt-guided-task-aware-image","slug":"promptcap-prompt-guided-task-aware-image","title":"PromptCap: Prompt-Guided Task-Aware Image Captioning","date":"2022-11-15","arxiv_id":"2211.09699","n_code_links":1,"syntology":null},{"paper":null,"slug":"robbert-2022-updating-a-dutch-language-model","title":"RobBERT-2022: Updating a Dutch Language Model to Account for Evolving Language Use","date":"2022-11-15","arxiv_id":"2211.08192","n_code_links":0,"syntology":null},{"paper":"/paper/are-hard-examples-also-harder-to-explain-a","slug":"are-hard-examples-also-harder-to-explain-a","title":"Are Hard Examples also Harder to Explain? A Study with Human and Model-Generated Explanations","date":"2022-11-14","arxiv_id":"2211.07517","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["swarnahub/explanationhardness"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/ugif-ui-grounded-instruction-following","slug":"ugif-ui-grounded-instruction-following","title":"UGIF: UI Grounded Instruction Following","date":"2022-11-14","arxiv_id":"2211.07615","n_code_links":0,"syntology":null},{"paper":null,"slug":"textual-data-augmentation-for-patient","title":"Textual Data Augmentation for Patient Outcomes Prediction","date":"2022-11-13","arxiv_id":"2211.06778","n_code_links":0,"syntology":null},{"paper":"/paper/what-would-harry-say-building-dialogue-agents","slug":"what-would-harry-say-building-dialogue-agents","title":"Large Language Models Meet Harry Potter: A Bilingual Dataset for Aligning Dialogue Agents with Characters","date":"2022-11-13","arxiv_id":"2211.06869","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-optimizing-the-communication-of-model","title":"On Optimizing the Communication of Model Parallelism","date":"2022-11-10","arxiv_id":"2211.05322","n_code_links":0,"syntology":null},{"paper":"/paper/collateral-facilitation-in-humans-and","slug":"collateral-facilitation-in-humans-and","title":"Collateral facilitation in humans and language models","date":"2022-11-09","arxiv_id":"2211.05198","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["jmichaelov/collateral-facilitation"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/active-example-selection-for-in-context","slug":"active-example-selection-for-in-context","title":"Active Example Selection for In-Context Learning","date":"2022-11-08","arxiv_id":"2211.04486","n_code_links":1,"syntology":{"ran":10,"of":16,"n_ran_checked":4,"n_instrument":6,"unverified":6,"pointer_only":0,"phrase":"10 ran (of which 3 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 6 where Syntology's instrument failed) · 6 unverified","official":{"repos":["chicagohai/active-example-selection"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":3,"n_ran_no_instrument_failure":4,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":"/paper/towards-asteroid-detection-in-microlensing","slug":"towards-asteroid-detection-in-microlensing","title":"Towards Asteroid Detection in Microlensing Surveys with Deep Learning","date":"2022-11-04","arxiv_id":"2211.02239","n_code_links":1,"syntology":null},{"paper":null,"slug":"using-large-pre-trained-language-model-to","title":"Using Large Pre-Trained Language Model to Assist FDA in Premarket Medical Device","date":"2022-11-03","arxiv_id":"2212.01217","n_code_links":0,"syntology":null},{"paper":"/paper/interpretability-in-the-wild-a-circuit-for","slug":"interpretability-in-the-wild-a-circuit-for","title":"Interpretability in the Wild: a Circuit for Indirect Object Identification in GPT-2 small","date":"2022-11-01","arxiv_id":"2211.00593","n_code_links":7,"syntology":{"ran":9,"of":13,"n_ran_checked":9,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["redwoodresearch/easy-transformer"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/text-only-training-for-image-captioning-using","slug":"text-only-training-for-image-captioning-using","title":"Text-Only Training for Image Captioning using Noise-Injected CLIP","date":"2022-11-01","arxiv_id":"2211.00575","n_code_links":4,"syntology":{"ran":4,"of":5,"n_ran_checked":2,"n_instrument":2,"unverified":1,"pointer_only":1,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","official":{"repos":["davidhuji/capdec"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/gptq-accurate-post-training-quantization-for","slug":"gptq-accurate-post-training-quantization-for","title":"GPTQ: Accurate Post-Training Quantization for Generative Pre-trained Transformers","date":"2022-10-31","arxiv_id":"2210.17323","n_code_links":17,"syntology":{"ran":5,"of":15,"n_ran_checked":2,"n_instrument":3,"unverified":10,"pointer_only":1,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 10 unverified","official":{"repos":["ist-daslab/gptq"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/ssd-lm-semi-autoregressive-simplex-based","slug":"ssd-lm-semi-autoregressive-simplex-based","title":"SSD-LM: Semi-autoregressive Simplex-based Diffusion Language Model for Text Generation and Modular Control","date":"2022-10-31","arxiv_id":"2210.17432","n_code_links":2,"syntology":null},{"paper":null,"slug":"towards-zero-shot-and-few-shot-table-question","title":"Towards Zero-Shot and Few-Shot Table Question Answering using GPT-3","date":"2022-10-31","arxiv_id":"2210.17284","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-decompose-hypothetical-question","title":"Learning to Decompose: Hypothetical Question Decomposition Based on Comparable Texts","date":"2022-10-30","arxiv_id":"2210.16865","n_code_links":0,"syntology":null},{"paper":"/paper/probing-for-targeted-syntactic-knowledge","slug":"probing-for-targeted-syntactic-knowledge","title":"Probing for targeted syntactic knowledge through grammatical error detection","date":"2022-10-28","arxiv_id":"2210.16228","n_code_links":1,"syntology":null},{"paper":null,"slug":"roma-run-time-object-detection-to-maximize","title":"ROMA: Run-Time Object Detection To Maximize Real-Time Accuracy","date":"2022-10-28","arxiv_id":"2210.16083","n_code_links":0,"syntology":null},{"paper":"/paper/coco-dr-combating-distribution-shifts-in-zero","slug":"coco-dr-combating-distribution-shifts-in-zero","title":"COCO-DR: Combating Distribution Shifts in Zero-Shot Dense Retrieval with Contrastive and Distributionally Robust Learning","date":"2022-10-27","arxiv_id":"2210.15212","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; the one sample that ran constructed an object rather than computing a result","official":{"repos":["openmatch/coco-dr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"trscore-a-novel-gpt-based-readability-scorer","title":"TRScore: A Novel GPT-based Readability Scorer for ASR Segmentation and Punctuation model evaluation and selection","date":"2022-10-27","arxiv_id":"2210.15104","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-robustness-of-prefix-tuning-in","title":"Exploring Robustness of Prefix Tuning in Noisy Data: A Case Study in Financial Sentiment Analysis","date":"2022-10-26","arxiv_id":"2211.05584","n_code_links":0,"syntology":null},{"paper":null,"slug":"ielm-an-open-information-extraction-benchmark","title":"IELM: An Open Information Extraction Benchmark for Pre-Trained Language Models","date":"2022-10-25","arxiv_id":"2210.14128","n_code_links":0,"syntology":null},{"paper":null,"slug":"xricl-cross-lingual-retrieval-augmented-in","title":"XRICL: Cross-lingual Retrieval-Augmented In-Context Learning for Cross-lingual Text-to-SQL Semantic Parsing","date":"2022-10-25","arxiv_id":"2210.13693","n_code_links":0,"syntology":null},{"paper":"/paper/abductive-action-inference","slug":"abductive-action-inference","title":"Inferring Past Human Actions in Homes with Abductive Reasoning","date":"2022-10-24","arxiv_id":"2210.13984","n_code_links":1,"syntology":null},{"paper":"/paper/emergent-world-representations-exploring-a","slug":"emergent-world-representations-exploring-a","title":"Emergent World Representations: Exploring a Sequence Model Trained on a Synthetic Task","date":"2022-10-24","arxiv_id":"2210.13382","n_code_links":4,"syntology":{"ran":6,"of":6,"n_ran_checked":5,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["likenneth/othello_world"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/exploring-euphemism-detection-in-few-shot-and","slug":"exploring-euphemism-detection-in-few-shot-and","title":"Exploring Euphemism Detection in Few-Shot and Zero-Shot Settings","date":"2022-10-24","arxiv_id":"2210.12926","n_code_links":1,"syntology":null},{"paper":"/paper/perfectly-secure-steganography-using-minimum","slug":"perfectly-secure-steganography-using-minimum","title":"Perfectly Secure Steganography Using Minimum Entropy Coupling","date":"2022-10-24","arxiv_id":"2210.14889","n_code_links":2,"syntology":{"ran":1,"of":4,"n_ran_checked":1,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["schroederdewitt/perfectly-secure-steganography"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/leveraging-large-language-models-for-multiple","slug":"leveraging-large-language-models-for-multiple","title":"Leveraging Large Language Models for Multiple Choice Question Answering","date":"2022-10-22","arxiv_id":"2210.12353","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":3,"n_instrument":3,"unverified":1,"pointer_only":0,"phrase":"6 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["byu-pccl/leveraging-llms-for-mcqa"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"meta-learning-pathologies-from-radiology","title":"Meta-learning Pathologies from Radiology Reports using Variance Aware Prototypical Networks","date":"2022-10-22","arxiv_id":"2210.13979","n_code_links":0,"syntology":null},{"paper":"/paper/a-causal-framework-to-quantify-the-robustness","slug":"a-causal-framework-to-quantify-the-robustness","title":"A Causal Framework to Quantify the Robustness of Mathematical Reasoning with Language Models","date":"2022-10-21","arxiv_id":"2210.12023","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":3,"n_instrument":2,"unverified":2,"pointer_only":7,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["alestolfo/causal-math"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/diffuser-efficient-transformers-with-multi","slug":"diffuser-efficient-transformers-with-multi","title":"Diffuser: Efficient Transformers with Multi-hop Attention Diffusion for Long Sequences","date":"2022-10-21","arxiv_id":"2210.11794","n_code_links":1,"syntology":null},{"paper":null,"slug":"wikiwhy-answering-and-explaining-cause-and","title":"WikiWhy: Answering and Explaining Cause-and-Effect Questions","date":"2022-10-21","arxiv_id":"2210.12152","n_code_links":0,"syntology":null},{"paper":null,"slug":"3dall-e-integrating-text-to-image-ai-in-3d","title":"3DALL-E: Integrating Text-to-Image AI in 3D Design Workflows","date":"2022-10-20","arxiv_id":"2210.11603","n_code_links":0,"syntology":null},{"paper":"/paper/composing-ensembles-of-pre-trained-models-via","slug":"composing-ensembles-of-pre-trained-models-via","title":"Composing Ensembles of Pre-trained Models via Iterative Consensus","date":"2022-10-20","arxiv_id":"2210.11522","n_code_links":0,"syntology":null},{"paper":"/paper/general-image-descriptors-for-open-world","slug":"general-image-descriptors-for-open-world","title":"General Image Descriptors for Open World Image Retrieval using ViT CLIP","date":"2022-10-20","arxiv_id":"2210.11141","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ivanaer/g-universal-clip"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/biogpt-generative-pre-trained-transformer-for","slug":"biogpt-generative-pre-trained-transformer-for","title":"BioGPT: Generative Pre-trained Transformer for Biomedical Text Generation and Mining","date":"2022-10-19","arxiv_id":"2210.10341","n_code_links":4,"syntology":null},{"paper":null,"slug":"towards-a-neural-architecture-of-language","title":"Towards a neural architecture of language: Deep learning versus logistics of access in neural architectures for compositional processing","date":"2022-10-19","arxiv_id":"2210.10543","n_code_links":0,"syntology":null},{"paper":null,"slug":"systematicity-in-gpt-3-s-interpretation-of","title":"Systematicity in GPT-3's Interpretation of Novel English Noun Compounds","date":"2022-10-18","arxiv_id":"2210.09492","n_code_links":0,"syntology":null},{"paper":null,"slug":"team-flow-at-drc2022-pipeline-system-for","title":"Team Flow at DRC2022: Pipeline System for Travel Destination Recommendation Task in Spoken Dialogue","date":"2022-10-18","arxiv_id":"2210.09518","n_code_links":0,"syntology":null},{"paper":null,"slug":"tiny-attention-adapter-contexts-are-more","title":"Tiny-Attention Adapter: Contexts Are More Important Than the Number of Parameters","date":"2022-10-18","arxiv_id":"2211.01979","n_code_links":0,"syntology":null},{"paper":"/paper/a-generative-user-simulator-with-gpt-based","slug":"a-generative-user-simulator-with-gpt-based","title":"A Generative User Simulator with GPT-based Architecture and Goal State Tracking for Reinforced Multi-Domain Dialog Systems","date":"2022-10-17","arxiv_id":"2210.08692","n_code_links":1,"syntology":{"ran":8,"of":9,"n_ran_checked":8,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["thu-spmi/gus"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/prompting-gpt-3-to-be-reliable","slug":"prompting-gpt-3-to-be-reliable","title":"Prompting GPT-3 To Be Reliable","date":"2022-10-17","arxiv_id":"2210.09150","n_code_links":1,"syntology":null},{"paper":"/paper/normsage-multi-lingual-multi-cultural-norm","slug":"normsage-multi-lingual-multi-cultural-norm","title":"NormSAGE: Multi-Lingual Multi-Cultural Norm Discovery from Conversations On-the-Fly","date":"2022-10-16","arxiv_id":"2210.08604","n_code_links":1,"syntology":null},{"paper":null,"slug":"object-attentional-untargeted-adversarial","title":"Object-Attentional Untargeted Adversarial Attack","date":"2022-10-16","arxiv_id":"2210.08472","n_code_links":0,"syntology":null},{"paper":"/paper/dylora-parameter-efficient-tuning-of-pre","slug":"dylora-parameter-efficient-tuning-of-pre","title":"DyLoRA: Parameter Efficient Tuning of Pre-trained Models using Dynamic Search-Free Low-Rank Adaptation","date":"2022-10-14","arxiv_id":"2210.07558","n_code_links":2,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["huawei-noah/kd-nlp"],"state":"official: not harvested","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":[]}}},{"paper":"/paper/extracting-cultural-commonsense-knowledge-at","slug":"extracting-cultural-commonsense-knowledge-at","title":"Extracting Cultural Commonsense Knowledge at Scale","date":"2022-10-14","arxiv_id":"2210.07763","n_code_links":2,"syntology":null},{"paper":"/paper/john-is-50-years-old-can-his-son-be-65","slug":"john-is-50-years-old-can-his-son-be-65","title":"\"John is 50 years old, can his son be 65?\" Evaluating NLP Models' Understanding of Feasibility","date":"2022-10-14","arxiv_id":"2210.07471","n_code_links":1,"syntology":null},{"paper":"/paper/testaug-a-framework-for-augmenting-capability-1","slug":"testaug-a-framework-for-augmenting-capability-1","title":"TestAug: A Framework for Augmenting Capability-based NLP Tests","date":"2022-10-14","arxiv_id":"2210.08097","n_code_links":1,"syntology":null},{"paper":"/paper/constructing-natural-language-explanations","slug":"constructing-natural-language-explanations","title":"Saliency Map Verbalization: Comparing Feature Importance Representations from Model-free and Instruction-based Methods","date":"2022-10-13","arxiv_id":"2210.07222","n_code_links":1,"syntology":null},{"paper":null,"slug":"explanations-from-large-language-models-make","title":"Explanations from Large Language Models Make Small Reasoners Better","date":"2022-10-13","arxiv_id":"2210.06726","n_code_links":0,"syntology":null},{"paper":null,"slug":"jointly-reinforced-user-simulator-and-task-1","title":"Jointly Reinforced User Simulator and Task-oriented Dialog System with Simplified Generative Architecture","date":"2022-10-13","arxiv_id":"2210.06706","n_code_links":0,"syntology":null},{"paper":"/paper/language-models-of-code-are-few-shot","slug":"language-models-of-code-are-few-shot","title":"Language Models of Code are Few-Shot Commonsense Learners","date":"2022-10-13","arxiv_id":"2210.07128","n_code_links":2,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["madaan/cocogen"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/large-language-models-are-few-1-shot-table","slug":"large-language-models-are-few-1-shot-table","title":"Large Language Models are few(1)-shot Table Reasoners","date":"2022-10-13","arxiv_id":"2210.06710","n_code_links":1,"syntology":null},{"paper":null,"slug":"are-sample-efficient-nlp-models-more-robust","title":"Are Sample-Efficient NLP Models More Robust?","date":"2022-10-12","arxiv_id":"2210.06456","n_code_links":0,"syntology":null},{"paper":"/paper/foundation-transformers","slug":"foundation-transformers","title":"Foundation Transformers","date":"2022-10-12","arxiv_id":"2210.06423","n_code_links":4,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["microsoft/unilm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/predictive-querying-for-autoregressive-neural","slug":"predictive-querying-for-autoregressive-neural","title":"Predictive Querying for Autoregressive Neural Sequence Models","date":"2022-10-12","arxiv_id":"2210.06464","n_code_links":1,"syntology":{"ran":7,"of":13,"n_ran_checked":2,"n_instrument":5,"unverified":6,"pointer_only":0,"phrase":"7 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 5 where Syntology's instrument failed) · 6 unverified","official":{"repos":["ajboyd2/prob_seq_queries"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":2,"n_ran_no_instrument_failure":2,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"sumbot-summarizing-context-in-open-domain","title":"SUMBot: Summarizing Context in Open-Domain Dialogue Systems","date":"2022-10-12","arxiv_id":"2210.06496","n_code_links":0,"syntology":null},{"paper":"/paper/rev-information-theoretic-evaluation-of-free","slug":"rev-information-theoretic-evaluation-of-free","title":"REV: Information-Theoretic Evaluation of Free-Text Rationales","date":"2022-10-10","arxiv_id":"2210.04982","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["hanjiechen/rev"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"the-minimum-wage-as-an-anchor-effects-on","title":"The Minimum Wage as an Anchor: Effects on Determinations of Fairness by Humans and AI","date":"2022-10-10","arxiv_id":"2210.10585","n_code_links":0,"syntology":null},{"paper":"/paper/asdot-any-shot-data-to-text-generation-with","slug":"asdot-any-shot-data-to-text-generation-with","title":"ASDOT: Any-Shot Data-to-Text Generation with Pretrained Language Models","date":"2022-10-09","arxiv_id":"2210.04325","n_code_links":1,"syntology":null},{"paper":"/paper/controllable-dialogue-simulation-with-in","slug":"controllable-dialogue-simulation-with-in","title":"Controllable Dialogue Simulation with In-Context Learning","date":"2022-10-09","arxiv_id":"2210.04185","n_code_links":1,"syntology":null},{"paper":"/paper/fine-tuning-pre-trained-transformers-into","slug":"fine-tuning-pre-trained-transformers-into","title":"Fine-Tuning Pre-trained Transformers into Decaying Fast Weights","date":"2022-10-09","arxiv_id":"2210.04243","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["jenni-ai/t2fw"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"alphatuning-quantization-aware-parameter","title":"AlphaTuning: Quantization-Aware Parameter-Efficient Adaptation of Large-Scale Pre-Trained Language Models","date":"2022-10-08","arxiv_id":"2210.03858","n_code_links":0,"syntology":null},{"paper":"/paper/automatic-chain-of-thought-prompting-in-large","slug":"automatic-chain-of-thought-prompting-in-large","title":"Automatic Chain of Thought Prompting in Large Language Models","date":"2022-10-07","arxiv_id":"2210.03493","n_code_links":5,"syntology":{"ran":1,"of":3,"n_ran_checked":0,"n_instrument":1,"unverified":2,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["amazon-science/auto-cot"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["named_in_paper"]}}},{"paper":"/paper/how-large-language-models-are-transforming","slug":"how-large-language-models-are-transforming","title":"How Large Language Models are Transforming Machine-Paraphrased Plagiarism","date":"2022-10-07","arxiv_id":"2210.03568","n_code_links":3,"syntology":null},{"paper":"/paper/in-search-of-a-robust-facial-expressions","slug":"in-search-of-a-robust-facial-expressions","title":"In Search of a Robust Facial Expressions Recognition Model: A Large-Scale Visual Cross-Corpus Study","date":"2022-10-07","arxiv_id":null,"n_code_links":1,"syntology":null}],"record_sha256":"1cf75685f3a419a9e28356d4c7a82097910267c93f9109ad53af76ad756bbfe1","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}