{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/gpt-3/papers/18","list_of":"/method/gpt-3","method":"GPT-3","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":18,"pages_in_order":20,"rows_per_page":100,"rows":[1701,1800],"of":1906,"counts":{"archive_papers_tagged":1906,"with_a_code_link":866,"where_syntology_ran_a_sample":319,"not_listed_spam_title":0,"listed":1906,"listed_where_code_ran":319,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":259,"every_run_a_failure_of_syntologys_instrument":60,"listed_with_a_run_with_no_instrument_failure":259,"listed_every_run_a_failure_of_syntologys_instrument":60,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/gpt-3","prev":"/method/gpt-3/papers/17","next":"/method/gpt-3/papers/19","papers":[{"paper":null,"slug":"word-play-for-playing-othello-reverses","title":"Word Play for Playing Othello (Reverses)","date":"2022-07-18","arxiv_id":"2207.08766","n_code_links":0,"syntology":null},{"paper":"/paper/can-large-language-models-reason-about","slug":"can-large-language-models-reason-about","title":"Can large language models reason about medical questions?","date":"2022-07-17","arxiv_id":"2207.08143","n_code_links":1,"syntology":{"ran":10,"of":10,"n_ran_checked":10,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["vlievin/medical-reasoning"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/electra-is-a-zero-shot-learner-too","slug":"electra-is-a-zero-shot-learner-too","title":"ELECTRA is a Zero-Shot Learner, Too","date":"2022-07-17","arxiv_id":"2207.08141","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nishiwen1214/rtd-electra"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/re2g-retrieve-rerank-generate-2","slug":"re2g-retrieve-rerank-generate-2","title":"Re2G: Retrieve, Rerank, Generate","date":"2022-07-13","arxiv_id":"2207.06300","n_code_links":1,"syntology":{"ran":7,"of":8,"n_ran_checked":6,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ibm/kgi-slot-filling"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"few-shot-training-llms-for-project-specific","title":"Few-shot training LLMs for project-specific code-summarization","date":"2022-07-09","arxiv_id":"2207.04237","n_code_links":0,"syntology":null},{"paper":null,"slug":"ask-me-what-you-need-product-retrieval-using","title":"Ask Me What You Need: Product Retrieval using Knowledge from GPT-3","date":"2022-07-06","arxiv_id":"2207.02516","n_code_links":0,"syntology":null},{"paper":null,"slug":"machine-learning-model-sizes-and-the","title":"Machine Learning Model Sizes and the Parameter Gap","date":"2022-07-05","arxiv_id":"2207.02852","n_code_links":0,"syntology":null},{"paper":"/paper/gaitforemer-self-supervised-pre-training-of","slug":"gaitforemer-self-supervised-pre-training-of","title":"GaitForeMer: Self-Supervised Pre-Training of Transformers via Human Motion Forecasting for Few-Shot Gait Impairment Severity Estimation","date":"2022-06-30","arxiv_id":"2207.00106","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-test-for-evaluating-performance-in-human","title":"A Test for Evaluating Performance in Human-Computer Systems","date":"2022-06-24","arxiv_id":"2206.12390","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-disability-lens-towards-biases-in-gpt-3","title":"A Disability Lens towards Biases in GPT-3 Generated Open-Ended Languages","date":"2022-06-23","arxiv_id":"2206.11993","n_code_links":0,"syntology":null},{"paper":null,"slug":"using-cognitive-psychology-to-understand-gpt","title":"Using cognitive psychology to understand GPT-3","date":"2022-06-21","arxiv_id":"2206.14576","n_code_links":0,"syntology":null},{"paper":"/paper/nuqmm-quantized-matmul-for-efficient","slug":"nuqmm-quantized-matmul-for-efficient","title":"LUT-GEMM: Quantized Matrix Multiplication based on LUTs for Efficient Inference in Large-Scale Generative Language Models","date":"2022-06-20","arxiv_id":"2206.09557","n_code_links":2,"syntology":null},{"paper":"/paper/argumentative-text-generation-in-economic","slug":"argumentative-text-generation-in-economic","title":"Argumentative Text Generation in Economic Domain","date":"2022-06-18","arxiv_id":"2206.09251","n_code_links":1,"syntology":null},{"paper":null,"slug":"automatic-summarization-of-russian-texts","title":"Automatic Summarization of Russian Texts: Comparison of Extractive and Abstractive Methods","date":"2022-06-18","arxiv_id":"2206.09253","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-dataset-and-benchmark-for-automatically","title":"From Human Days to Machine Seconds: Automatically Answering and Generating Machine Learning Final Exams","date":"2022-06-11","arxiv_id":"2206.05442","n_code_links":0,"syntology":null},{"paper":"/paper/putting-gpt-3-s-creativity-to-the-alternative","slug":"putting-gpt-3-s-creativity-to-the-alternative","title":"Putting GPT-3's Creativity to the (Alternative Uses) Test","date":"2022-06-10","arxiv_id":"2206.08932","n_code_links":1,"syntology":null},{"paper":null,"slug":"dynamar-dynamic-prompt-with-mask-token","title":"DynaMaR: Dynamic Prompt with Mask Token Representation","date":"2022-06-07","arxiv_id":"2206.02982","n_code_links":0,"syntology":null},{"paper":"/paper/on-the-advance-of-making-language-models","slug":"on-the-advance-of-making-language-models","title":"Making Large Language Models Better Reasoners with Step-Aware Verifier","date":"2022-06-06","arxiv_id":"2206.02336","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-generation-of-programming-exercises-1","title":"Automatic Generation of Programming Exercises and Code Explanations using Large Language Models","date":"2022-06-03","arxiv_id":"2206.11861","n_code_links":0,"syntology":null},{"paper":"/paper/decentralized-training-of-foundation-models","slug":"decentralized-training-of-foundation-models","title":"Decentralized Training of Foundation Models in Heterogeneous Environments","date":"2022-06-02","arxiv_id":"2206.01288","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":1,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["DS3Lab/DT-FM"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"romantic-computing","title":"Romantic-Computing","date":"2022-06-01","arxiv_id":"2206.11864","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-graph-deep-learning-a-case-study-in","title":"Knowledge Graph - Deep Learning: A Case Study in Question Answering in Aviation Safety Domain","date":"2022-05-31","arxiv_id":"2205.15952","n_code_links":0,"syntology":null},{"paper":"/paper/billions-of-parameters-are-worth-more-than-in","slug":"billions-of-parameters-are-worth-more-than-in","title":"Billions of Parameters Are Worth More Than In-domain Training Data: A case study in the Legal Case Entailment Task","date":"2022-05-30","arxiv_id":"2205.15172","n_code_links":1,"syntology":null},{"paper":"/paper/teaching-models-to-express-their-uncertainty","slug":"teaching-models-to-express-their-uncertainty","title":"Teaching Models to Express Their Uncertainty in Words","date":"2022-05-28","arxiv_id":"2205.14334","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 2 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["sylinrl/calibratedmath"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"conditional-set-generation-using-seq2seq-1","title":"Conditional set generation using Seq2seq models","date":"2022-05-25","arxiv_id":"2205.12485","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-are-zero-shot-clinical","title":"Large Language Models are Few-Shot Clinical Information Extractors","date":"2022-05-25","arxiv_id":"2205.12689","n_code_links":0,"syntology":null},{"paper":"/paper/naturalprover-grounded-mathematical-proof","slug":"naturalprover-grounded-mathematical-proof","title":"NaturalProver: Grounded Mathematical Proof Generation with Language Models","date":"2022-05-25","arxiv_id":"2205.12910","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["wellecks/naturalprover"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/flute-figurative-language-understanding-and","slug":"flute-figurative-language-understanding-and","title":"FLUTE: Figurative Language Understanding through Textual Explanations","date":"2022-05-24","arxiv_id":"2205.12404","n_code_links":1,"syntology":null},{"paper":null,"slug":"how-human-is-human-evaluation-improving-the","title":"The Authenticity Gap in Human Evaluation","date":"2022-05-24","arxiv_id":"2205.11930","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-short-text-classification-with","title":"Improving Short Text Classification With Augmented Data Using GPT-3","date":"2022-05-23","arxiv_id":"2205.10981","n_code_links":0,"syntology":null},{"paper":"/paper/looking-for-a-handsome-carpenter-debiasing","slug":"looking-for-a-handsome-carpenter-debiasing","title":"Looking for a Handsome Carpenter! Debiasing GPT-3 Job Advertisements","date":"2022-05-23","arxiv_id":"2205.11374","n_code_links":1,"syntology":null},{"paper":null,"slug":"penguins-don-t-fly-reasoning-about-generics","title":"Penguins Don't Fly: Reasoning about Generics through Instantiations and Exceptions","date":"2022-05-23","arxiv_id":"2205.11658","n_code_links":0,"syntology":null},{"paper":null,"slug":"rl-with-kl-penalties-is-better-viewed-as","title":"RL with KL penalties is better viewed as Bayesian inference","date":"2022-05-23","arxiv_id":"2205.11275","n_code_links":0,"syntology":null},{"paper":"/paper/instruction-induction-from-few-examples-to","slug":"instruction-induction-from-few-examples-to","title":"Instruction Induction: From Few Examples to Natural Language Task Descriptions","date":"2022-05-22","arxiv_id":"2205.10782","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["orhonovich/instruction-induction"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/least-to-most-prompting-enables-complex","slug":"least-to-most-prompting-enables-complex","title":"Least-to-Most Prompting Enables Complex Reasoning in Large Language Models","date":"2022-05-21","arxiv_id":"2205.10625","n_code_links":1,"syntology":null},{"paper":null,"slug":"m6-rec-generative-pretrained-language-models","title":"M6-Rec: Generative Pretrained Language Models are Open-Ended Recommender Systems","date":"2022-05-17","arxiv_id":"2205.08084","n_code_links":0,"syntology":null},{"paper":null,"slug":"heroes-villains-and-victims-and-gpt-3","title":"Heroes, Villains, and Victims, and GPT-3: Automated Extraction of Character Roles Without Training Data","date":"2022-05-16","arxiv_id":"2205.07557","n_code_links":0,"syntology":null},{"paper":"/paper/the-ai-teacher-test-measuring-the-pedagogical","slug":"the-ai-teacher-test-measuring-the-pedagogical","title":"The AI Teacher Test: Measuring the Pedagogical Ability of Blender and GPT-3 in Educational Dialogues","date":"2022-05-16","arxiv_id":"2205.07540","n_code_links":1,"syntology":null},{"paper":"/paper/clinical-prompt-learning-with-frozen-language","slug":"clinical-prompt-learning-with-frozen-language","title":"Clinical Prompt Learning with Frozen Language Models","date":"2022-05-11","arxiv_id":"2205.05535","n_code_links":1,"syntology":null},{"paper":"/paper/towards-the-generation-of-musical","slug":"towards-the-generation-of-musical","title":"Towards the Generation of Musical Explanations with GPT-3","date":"2022-05-11","arxiv_id":"2206.08264","n_code_links":1,"syntology":null},{"paper":"/paper/reducing-activation-recomputation-in-large","slug":"reducing-activation-recomputation-in-large","title":"Reducing Activation Recomputation in Large Transformer Models","date":"2022-05-10","arxiv_id":"2205.05198","n_code_links":4,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["NVIDIA/Megatron-LM"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["listed"]}}},{"paper":"/paper/unifying-language-learning-paradigms","slug":"unifying-language-learning-paradigms","title":"UL2: Unifying Language Learning Paradigms","date":"2022-05-10","arxiv_id":"2205.05131","n_code_links":2,"syntology":{"ran":15,"of":16,"n_ran_checked":15,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"15 ran (of which 0 constructed an object rather than computing a result; 15 with no instrument failure: 0 honoured, 0 violated, 15 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["google-research/google-research"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/the-unreliability-of-explanations-in-few-shot","slug":"the-unreliability-of-explanations-in-few-shot","title":"The Unreliability of Explanations in Few-shot Prompting for Textual Reasoning","date":"2022-05-06","arxiv_id":"2205.03401","n_code_links":1,"syntology":null},{"paper":"/paper/when-a-sentence-does-not-introduce-a-1","slug":"when-a-sentence-does-not-introduce-a-1","title":"When a sentence does not introduce a discourse entity, Transformer-based models still sometimes refer to it","date":"2022-05-06","arxiv_id":"2205.03472","n_code_links":1,"syntology":null},{"paper":"/paper/contrastive-learning-for-prompt-based-few","slug":"contrastive-learning-for-prompt-based-few","title":"Contrastive Learning for Prompt-Based Few-Shot Language Learners","date":"2022-05-03","arxiv_id":"2205.01308","n_code_links":1,"syntology":{"ran":6,"of":11,"n_ran_checked":4,"n_instrument":2,"unverified":5,"pointer_only":6,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 1 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 5 unverified","official":{"repos":["yiren-jian/lm-supcon"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/mixed-effects-transformers-for-hierarchical","slug":"mixed-effects-transformers-for-hierarchical","title":"Mixed-effects transformers for hierarchical adaptation","date":"2022-05-03","arxiv_id":"2205.01749","n_code_links":1,"syntology":null},{"paper":"/paper/opt-open-pre-trained-transformer-language","slug":"opt-open-pre-trained-transformer-language","title":"OPT: Open Pre-trained Transformer Language Models","date":"2022-05-02","arxiv_id":"2205.01068","n_code_links":11,"syntology":{"ran":14,"of":24,"n_ran_checked":14,"n_instrument":0,"unverified":10,"pointer_only":17,"phrase":"14 ran (of which 2 constructed an object rather than computing a result; 14 with no instrument failure: 0 honoured, 0 violated, 14 with no contract checked; 0 where Syntology's instrument failed) · 10 unverified","official":{"repos":["facebookresearch/metaseq"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":3,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"learning-from-natural-language-feedback","title":"Training Language Models with Language Feedback","date":"2022-04-29","arxiv_id":"2204.14146","n_code_links":0,"syntology":null},{"paper":"/paper/inferring-implicit-relations-with-language","slug":"inferring-implicit-relations-with-language","title":"Inferring Implicit Relations in Complex Questions with Language Models","date":"2022-04-28","arxiv_id":"2204.13778","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-the-effect-of-pretraining-corpora-on-in","title":"On the Effect of Pretraining Corpora on In-context Learning by a Large-scale Language Model","date":"2022-04-28","arxiv_id":"2204.13509","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-end-to-end-dialogue-summarization-system","title":"An End-to-End Dialogue Summarization System for Sales Calls","date":"2022-04-27","arxiv_id":"2204.12951","n_code_links":0,"syntology":null},{"paper":null,"slug":"codexdb-generating-code-for-processing-sql","title":"CodexDB: Generating Code for Processing SQL Queries using GPT-3 Codex","date":"2022-04-19","arxiv_id":"2204.08941","n_code_links":0,"syntology":null},{"paper":"/paper/mgpt-few-shot-learners-go-multilingual","slug":"mgpt-few-shot-learners-go-multilingual","title":"mGPT: Few-Shot Learners Go Multilingual","date":"2022-04-15","arxiv_id":"2204.07580","n_code_links":1,"syntology":null},{"paper":"/paper/gpt-neox-20b-an-open-source-autoregressive-1","slug":"gpt-neox-20b-an-open-source-autoregressive-1","title":"GPT-NeoX-20B: An Open-Source Autoregressive Language Model","date":"2022-04-14","arxiv_id":"2204.06745","n_code_links":11,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["eleutherai/gpt-neox"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"rows-from-many-sources-enriching-row","title":"Rows from Many Sources: Enriching row completions from Wikidata with a pre-trained Language Model","date":"2022-04-14","arxiv_id":"2204.07014","n_code_links":0,"syntology":null},{"paper":null,"slug":"bertuit-understanding-spanish-language-in","title":"BERTuit: Understanding Spanish language in Twitter through a native transformer","date":"2022-04-07","arxiv_id":"2204.03465","n_code_links":0,"syntology":null},{"paper":"/paper/data-augmentation-for-intent-classification-1","slug":"data-augmentation-for-intent-classification-1","title":"Data Augmentation for Intent Classification with Off-the-shelf Large Language Models","date":"2022-04-05","arxiv_id":"2204.01959","n_code_links":1,"syntology":{"ran":7,"of":10,"n_ran_checked":6,"n_instrument":1,"unverified":3,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["elementai/data-augmentation-with-llms"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"generative-pre-trained-transformers-for","title":"Generative Pre-Trained Transformers for Biologically Inspired Design","date":"2022-03-31","arxiv_id":"2204.09714","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-pre-trained-language-models-for","title":"Leveraging pre-trained language models for conversational information seeking from text","date":"2022-03-31","arxiv_id":"2204.03542","n_code_links":0,"syntology":null},{"paper":"/paper/transformer-language-models-without","slug":"transformer-language-models-without","title":"Transformer Language Models without Positional Encodings Still Learn Positional Information","date":"2022-03-30","arxiv_id":"2203.16634","n_code_links":1,"syntology":{"ran":8,"of":11,"n_ran_checked":8,"n_instrument":0,"unverified":3,"pointer_only":11,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["adihaviv/nopos"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/training-compute-optimal-large-language","slug":"training-compute-optimal-large-language","title":"Training Compute-Optimal Large Language Models","date":"2022-03-29","arxiv_id":"2203.15556","n_code_links":2,"syntology":{"ran":8,"of":11,"n_ran_checked":5,"n_instrument":3,"unverified":3,"pointer_only":4,"phrase":"8 ran (of which 3 constructed an object rather than computing a result; 5 with no instrument failure: 2 honoured, 0 violated, 3 with no contract checked; 3 where Syntology's instrument failed) · 3 unverified","official":null}},{"paper":"/paper/thinking-about-gpt-3-in-context-learning-for","slug":"thinking-about-gpt-3-in-context-learning-for","title":"Thinking about GPT-3 In-Context Learning for Biomedical IE? Think Again","date":"2022-03-16","arxiv_id":"2203.08410","n_code_links":1,"syntology":null},{"paper":null,"slug":"the-ghost-in-the-machine-has-an-american","title":"The Ghost in the Machine has an American accent: value conflict in GPT-3","date":"2022-03-15","arxiv_id":"2203.07785","n_code_links":0,"syntology":null},{"paper":"/paper/tensor-programs-v-tuning-large-neural","slug":"tensor-programs-v-tuning-large-neural","title":"Tensor Programs V: Tuning Large Neural Networks via Zero-Shot Hyperparameter Transfer","date":"2022-03-07","arxiv_id":"2203.03466","n_code_links":7,"syntology":{"ran":3,"of":5,"n_ran_checked":1,"n_instrument":2,"unverified":2,"pointer_only":0,"phrase":"3 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["microsoft/mup"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/training-language-models-to-follow","slug":"training-language-models-to-follow","title":"Training language models to follow instructions with human feedback","date":"2022-03-04","arxiv_id":"2203.02155","n_code_links":11,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["openai/following-instructions-human-feedback"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/rethinking-the-role-of-demonstrations-what","slug":"rethinking-the-role-of-demonstrations-what","title":"Rethinking the Role of Demonstrations: What Makes In-Context Learning Work?","date":"2022-02-25","arxiv_id":"2202.12837","n_code_links":2,"syntology":null},{"paper":"/paper/from-natural-language-to-simulations-applying","slug":"from-natural-language-to-simulations-applying","title":"From Natural Language to Simulations: Applying GPT-3 Codex to Automate Simulation Modeling of Logistics Systems","date":"2022-02-24","arxiv_id":"2202.12107","n_code_links":1,"syntology":null},{"paper":"/paper/semantic-features-of-object-concepts","slug":"semantic-features-of-object-concepts","title":"Semantic features of object concepts generated with GPT-3","date":"2022-02-08","arxiv_id":"2202.03753","n_code_links":1,"syntology":null},{"paper":"/paper/cedille-a-large-autoregressive-french","slug":"cedille-a-large-autoregressive-french","title":"Cedille: A large autoregressive French language model","date":"2022-02-07","arxiv_id":"2202.03371","n_code_links":1,"syntology":null},{"paper":null,"slug":"ethics-rules-of-engagement-and-ai-neural","title":"Ethics, Rules of Engagement, and AI: Neural Narrative Mapping Using Large Transformer Language Models","date":"2022-02-05","arxiv_id":"2202.02647","n_code_links":0,"syntology":null},{"paper":"/paper/co-training-improves-prompt-based-learning","slug":"co-training-improves-prompt-based-learning","title":"Co-training Improves Prompt-based Learning for Large Language Models","date":"2022-02-02","arxiv_id":"2202.00828","n_code_links":1,"syntology":{"ran":5,"of":10,"n_ran_checked":5,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["clinicalml/cotrain-prompting"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":"/paper/chain-of-thought-prompting-elicits-reasoning","slug":"chain-of-thought-prompting-elicits-reasoning","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","date":"2022-01-28","arxiv_id":"2201.11903","n_code_links":19,"syntology":{"ran":5,"of":7,"n_ran_checked":4,"n_instrument":1,"unverified":2,"pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":"/paper/summarizing-differences-between-text","slug":"summarizing-differences-between-text","title":"Describing Differences between Text Distributions with Natural Language","date":"2022-01-28","arxiv_id":"2201.12323","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":0,"n_instrument":4,"unverified":0,"pointer_only":4,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 4 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ruiqi-zhong/describedistributionaldifferences"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/synchromesh-reliable-code-generation-from-pre-1","slug":"synchromesh-reliable-code-generation-from-pre-1","title":"Synchromesh: Reliable code generation from pre-trained language models","date":"2022-01-26","arxiv_id":"2201.11227","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"whose-language-counts-as-high-quality","title":"Whose Language Counts as High Quality? Measuring Language Ideologies in Text Data Selection","date":"2022-01-25","arxiv_id":"2201.10474","n_code_links":0,"syntology":null},{"paper":null,"slug":"synthetic-books","title":"Synthetic Books","date":"2022-01-24","arxiv_id":"2201.09518","n_code_links":0,"syntology":null},{"paper":"/paper/black-box-prompt-learning-for-pre-trained","slug":"black-box-prompt-learning-for-pre-trained","title":"Black-box Prompt Learning for Pre-trained Language Models","date":"2022-01-21","arxiv_id":"2201.08531","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["shizhediao/black-box-prompt-learning"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/coauthor-designing-a-human-ai-collaborative","slug":"coauthor-designing-a-human-ai-collaborative","title":"CoAuthor: Designing a Human-AI Collaborative Writing Dataset for Exploring Language Model Capabilities","date":"2022-01-18","arxiv_id":"2201.06796","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-study-of-pre-trained-language-models-for","title":"A Study of Pre-trained Language Models for Analogy Generation","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"few-shot-semantic-parsing-with-language-1","title":"Few-Shot Semantic Parsing with Language Models Trained On Code","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"hierarchical-transformers-are-more-efficient-1","title":"Hierarchical Transformers Are More Efficient Language Models","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/memory-assisted-prompt-editing-to-improve-gpt","slug":"memory-assisted-prompt-editing-to-improve-gpt","title":"Memory-assisted prompt editing to improve GPT-3 after deployment","date":"2022-01-16","arxiv_id":"2201.06009","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["madaan/memprompt"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"penguins-dont-fly-reasoning-about-generics","title":"Penguins Don’t Fly: Reasoning about Generics through Instantiations and Exceptions","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"re2g-retrieve-rerank-generate","title":"Re2G: Retrieve, Rerank, Generate","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"reframing-human-ai-collaboration-for-1","title":"Reframing Human-AI Collaboration for Generating Free-Text Explanations","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/unifiedskg-unifying-and-multi-tasking","slug":"unifiedskg-unifying-and-multi-tasking","title":"UnifiedSKG: Unifying and Multi-Tasking Structured Knowledge Grounding with Text-to-Text Language Models","date":"2022-01-16","arxiv_id":"2201.05966","n_code_links":1,"syntology":{"ran":2,"of":5,"n_ran_checked":1,"n_instrument":1,"unverified":3,"pointer_only":1,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["hkunlp/unifiedskg"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/wanli-worker-and-ai-collaboration-for-natural","slug":"wanli-worker-and-ai-collaboration-for-natural","title":"WANLI: Worker and AI Collaboration for Natural Language Inference Dataset Creation","date":"2022-01-16","arxiv_id":"2201.05955","n_code_links":1,"syntology":null},{"paper":null,"slug":"when-a-sentence-does-not-introduce-a","title":"When a sentence does not introduce a discourse entity, Transformer-based models still often refer to it","date":"2022-01-16","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"commonsenseqa-2-0-exposing-the-limits-of-ai","title":"CommonsenseQA 2.0: Exposing the Limits of AI through Gamification","date":"2022-01-14","arxiv_id":"2201.05320","n_code_links":0,"syntology":null},{"paper":"/paper/black-box-tuning-for-language-model-as-a","slug":"black-box-tuning-for-language-model-as-a","title":"Black-Box Tuning for Language-Model-as-a-Service","date":"2022-01-10","arxiv_id":"2201.03514","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":3,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["txsun1997/black-box-tuning"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"computational-lens-on-cognition-study-of","title":"Imagined versus Remembered Stories: Quantifying Differences in Narrative Flow","date":"2022-01-07","arxiv_id":"2201.02662","n_code_links":0,"syntology":null},{"paper":"/paper/a-neural-network-solves-and-generates","slug":"a-neural-network-solves-and-generates","title":"A Neural Network Solves, Explains, and Generates University Math Problems by Program Synthesis and Few-Shot Learning at Human Level","date":"2021-12-31","arxiv_id":"2112.15594","n_code_links":1,"syntology":{"ran":4,"of":4,"n_ran_checked":4,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["idrori/mathq"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/ernie-3-0-titan-exploring-larger-scale","slug":"ernie-3-0-titan-exploring-larger-scale","title":"ERNIE 3.0 Titan: Exploring Larger-scale Knowledge Enhanced Pre-training for Language Understanding and Generation","date":"2021-12-23","arxiv_id":"2112.12731","n_code_links":3,"syntology":null},{"paper":"/paper/few-shot-learning-with-multilingual-language","slug":"few-shot-learning-with-multilingual-language","title":"Few-shot Learning with Multilingual Language Models","date":"2021-12-20","arxiv_id":"2112.10668","n_code_links":2,"syntology":null},{"paper":null,"slug":"analysis-and-mitigation-of-dataset-artifacts","title":"Analysis and Mitigation of Dataset Artifacts in OpenAI GPT-3","date":"2021-12-19","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/webgpt-browser-assisted-question-answering","slug":"webgpt-browser-assisted-question-answering","title":"WebGPT: Browser-assisted question-answering with human feedback","date":"2021-12-17","arxiv_id":"2112.09332","n_code_links":2,"syntology":null},{"paper":null,"slug":"few-shot-semantic-parsing-with-language","title":"Few-Shot Semantic Parsing with Language Models Trained On Code","date":"2021-12-16","arxiv_id":"2112.08696","n_code_links":0,"syntology":null},{"paper":"/paper/reframing-human-ai-collaboration-for","slug":"reframing-human-ai-collaboration-for","title":"Reframing Human-AI Collaboration for Generating Free-Text Explanations","date":"2021-12-16","arxiv_id":"2112.08674","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["allenai/few_shot_explanations"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/glam-efficient-scaling-of-language-models","slug":"glam-efficient-scaling-of-language-models","title":"GLaM: Efficient Scaling of Language Models with Mixture-of-Experts","date":"2021-12-13","arxiv_id":"2112.06905","n_code_links":0,"syntology":null},{"paper":"/paper/improving-language-models-by-retrieving-from","slug":"improving-language-models-by-retrieving-from","title":"Improving language models by retrieving from trillions of tokens","date":"2021-12-08","arxiv_id":"2112.04426","n_code_links":2,"syntology":{"ran":16,"of":23,"n_ran_checked":14,"n_instrument":2,"unverified":7,"pointer_only":3,"phrase":"16 ran (of which 5 constructed an object rather than computing a result; 14 with no instrument failure: 0 honoured, 3 violated, 11 with no contract checked; 2 where Syntology's instrument failed) · 7 unverified","official":null}}],"record_sha256":"da8951165169260c4b987a4ba3812b406011e8aae579668907b1eb5abb473f4c","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}