{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/gpt/papers/11","list_of":"/method/gpt","method":"GPT","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":11,"pages_in_order":13,"rows_per_page":100,"rows":[1001,1100],"of":1212,"counts":{"archive_papers_tagged":1212,"with_a_code_link":453,"where_syntology_ran_a_sample":152,"not_listed_spam_title":0,"listed":1212,"listed_where_code_ran":152,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":130,"every_run_a_failure_of_syntologys_instrument":22,"listed_with_a_run_with_no_instrument_failure":130,"listed_every_run_a_failure_of_syntologys_instrument":22,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/gpt","prev":"/method/gpt/papers/10","next":"/method/gpt/papers/12","papers":[{"paper":"/paper/medalpaca-an-open-source-collection-of","slug":"medalpaca-an-open-source-collection-of","title":"MedAlpaca -- An Open-Source Collection of Medical Conversational AI Models and Training Data","date":"2023-04-14","arxiv_id":"2304.08247","n_code_links":1,"syntology":{"ran":10,"of":13,"n_ran_checked":10,"n_instrument":0,"unverified":3,"pointer_only":1,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":null}},{"paper":null,"slug":"chatgpt-cites-the-most-cited-articles-and","title":"ChatGPT cites the most-cited articles and journals, relying solely on Google Scholar's citation counts. As a result, AI may amplify the Matthew Effect in environmental science","date":"2023-04-13","arxiv_id":"2304.06794","n_code_links":0,"syntology":null},{"paper":"/paper/shall-we-pretrain-autoregressive-language","slug":"shall-we-pretrain-autoregressive-language","title":"Shall We Pretrain Autoregressive Language Models with Retrieval? A Comprehensive Study","date":"2023-04-13","arxiv_id":"2304.06762","n_code_links":1,"syntology":null},{"paper":null,"slug":"approximating-human-evaluation-of-social","title":"Approximating Online Human Evaluation of Social Chatbots with Prompting","date":"2023-04-11","arxiv_id":"2304.05253","n_code_links":0,"syntology":null},{"paper":null,"slug":"distinguishing-chatgpt-3-5-4-generated-and","title":"Distinguishing ChatGPT(-3.5, -4)-generated and human-written papers through Japanese stylometric analysis","date":"2023-04-11","arxiv_id":"2304.05534","n_code_links":0,"syntology":null},{"paper":null,"slug":"training-large-language-models-efficiently","title":"Training Large Language Models Efficiently with Sparsity and Dataflow","date":"2023-04-11","arxiv_id":"2304.05511","n_code_links":0,"syntology":null},{"paper":"/paper/gpt-detectors-are-biased-against-non-native","slug":"gpt-detectors-are-biased-against-non-native","title":"GPT detectors are biased against non-native English writers","date":"2023-04-06","arxiv_id":"2304.02819","n_code_links":2,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["weixin-liang/chatgpt-detector-bias"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"geotechnical-parrot-tales-gpt-overcoming-gpt","title":"Geotechnical Parrot Tales (GPT): Harnessing Large Language Models in geotechnical engineering","date":"2023-04-04","arxiv_id":"2304.02138","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpt-4-to-gpt-3-5-hold-my-scalpel-a-look-at","title":"GPT-4 to GPT-3.5: 'Hold My Scalpel' -- A Look at the Competency of OpenAI's GPT on the Plastic Surgery In-Service Training Exam","date":"2023-04-04","arxiv_id":"2304.01503","n_code_links":0,"syntology":null},{"paper":null,"slug":"summary-of-chatgpt-gpt-4-research-and","title":"Summary of ChatGPT-Related Research and Perspective Towards the Future of Large Language Models","date":"2023-04-04","arxiv_id":"2304.01852","n_code_links":0,"syntology":null},{"paper":"/paper/greekbart-the-first-pretrained-greek-sequence","slug":"greekbart-the-first-pretrained-greek-sequence","title":"GreekBART: The First Pretrained Greek Sequence-to-Sequence Model","date":"2023-04-03","arxiv_id":"2304.00869","n_code_links":2,"syntology":null},{"paper":null,"slug":"aligning-a-medium-size-gpt-model-in-english","title":"Aligning a medium-size GPT model in English to a small closed domain in Spanish","date":"2023-03-30","arxiv_id":"2303.17649","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluation-of-gpt-and-bert-based-models-on","title":"Evaluation of GPT and BERT-based models on identifying protein-protein interactions in biomedical text","date":"2023-03-30","arxiv_id":"2303.17728","n_code_links":0,"syntology":null},{"paper":null,"slug":"humans-in-humans-out-on-gpt-converging-toward","title":"Humans in Humans Out: On GPT Converging Toward Common Sense in both Success and Failure","date":"2023-03-30","arxiv_id":"2303.17276","n_code_links":0,"syntology":null},{"paper":"/paper/autoad-movie-description-in-context","slug":"autoad-movie-description-in-context","title":"AutoAD: Movie Description in Context","date":"2023-03-29","arxiv_id":"2303.16899","n_code_links":1,"syntology":{"ran":5,"of":13,"n_ran_checked":3,"n_instrument":2,"unverified":8,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 8 unverified","official":{"repos":["Soldelli/MAD"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":8,"ran_from_kinds":["official"]}}},{"paper":"/paper/viewrefer-grasp-the-multi-view-knowledge-for","slug":"viewrefer-grasp-the-multi-view-knowledge-for","title":"ViewRefer: Grasp the Multi-view Knowledge for 3D Visual Grounding with GPT and Prototype Guidance","date":"2023-03-29","arxiv_id":"2303.16894","n_code_links":7,"syntology":{"ran":8,"of":12,"n_ran_checked":6,"n_instrument":2,"unverified":4,"pointer_only":12,"phrase":"8 ran (of which 2 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 1 violated, 5 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","official":{"repos":["ivan-tang-3d/viewrefer3d","ziyuguo99/viewrefer3d"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/zero-shot-clinical-entity-recognition-using","slug":"zero-shot-clinical-entity-recognition-using","title":"Improving Large Language Models for Clinical Named Entity Recognition via Prompt Engineering","date":"2023-03-29","arxiv_id":"2303.16416","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-analysis-of-gpt-3-s-performance-in","title":"Analyzing the Performance of GPT-3.5 and GPT-4 in Grammatical Error Correction","date":"2023-03-25","arxiv_id":"2303.14342","n_code_links":0,"syntology":null},{"paper":null,"slug":"gesgpt-speech-gesture-synthesis-with-text","title":"GesGPT: Speech Gesture Synthesis With Text Parsing from ChatGPT","date":"2023-03-23","arxiv_id":"2303.13013","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-complete-survey-on-generative-ai-aigc-is","title":"A Complete Survey on Generative AI (AIGC): Is ChatGPT from GPT-4 to GPT-5 All You Need?","date":"2023-03-21","arxiv_id":"2303.11717","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-comprehensive-capability-analysis-of-gpt-3","title":"A Comprehensive Capability Analysis of GPT-3 and GPT-3.5 Series Models","date":"2023-03-18","arxiv_id":"2303.10420","n_code_links":0,"syntology":null},{"paper":null,"slug":"spdf-sparse-pre-training-and-dense-fine","title":"SPDF: Sparse Pre-training and Dense Fine-tuning for Large Language Models","date":"2023-03-18","arxiv_id":"2303.10464","n_code_links":0,"syntology":null},{"paper":null,"slug":"gpts-are-gpts-an-early-look-at-the-labor","title":"GPTs are GPTs: An Early Look at the Labor Market Impact Potential of Large Language Models","date":"2023-03-17","arxiv_id":"2303.10130","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-generative-pre-trained-transformers-gpt","title":"Can Generative Pre-trained Transformers (GPT) Pass Assessments in Higher Education Programming Courses?","date":"2023-03-16","arxiv_id":"2303.09325","n_code_links":0,"syntology":null},{"paper":"/paper/evaluation-of-chatgpt-as-a-question-answering","slug":"evaluation-of-chatgpt-as-a-question-answering","title":"Can ChatGPT Replace Traditional KBQA Models? An In-depth Analysis of the Question Answering Performance of the GPT LLM Family","date":"2023-03-14","arxiv_id":"2303.07992","n_code_links":2,"syntology":null},{"paper":null,"slug":"algorithmic-ghost-in-the-research-shell-large","title":"Algorithmic Ghost in the Research Shell: Large Language Models and Academic Knowledge Creation in Management Research","date":"2023-03-10","arxiv_id":"2303.07304","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-gpt-struggle-to-answer","title":"Large Language Models (GPT) Struggle to Answer Multiple-Choice Questions about Code","date":"2023-03-09","arxiv_id":"2303.08033","n_code_links":0,"syntology":null},{"paper":"/paper/a-comprehensive-survey-of-ai-generated","slug":"a-comprehensive-survey-of-ai-generated","title":"A Comprehensive Survey of AI-Generated Content (AIGC): A History of Generative AI from GAN to ChatGPT","date":"2023-03-07","arxiv_id":"2303.04226","n_code_links":1,"syntology":null},{"paper":null,"slug":"cross-lingual-summarization-via-chatgpt","title":"Zero-Shot Cross-Lingual Summarization via Large Language Models","date":"2023-02-28","arxiv_id":"2302.14229","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-are-state-of-the-art","slug":"large-language-models-are-state-of-the-art","title":"Large Language Models Are State-of-the-Art Evaluators of Translation Quality","date":"2023-02-28","arxiv_id":"2302.14520","n_code_links":4,"syntology":null},{"paper":"/paper/large-scale-multi-modal-pre-trained-models-a","slug":"large-scale-multi-modal-pre-trained-models-a","title":"Large-scale Multi-Modal Pre-trained Models: A Comprehensive Survey","date":"2023-02-20","arxiv_id":"2302.10035","n_code_links":1,"syntology":null},{"paper":"/paper/how-good-are-gpt-models-at-machine","slug":"how-good-are-gpt-models-at-machine","title":"How Good Are GPT Models at Machine Translation? A Comprehensive Evaluation","date":"2023-02-18","arxiv_id":"2302.09210","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["microsoft/gpt-mt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/pac-prediction-sets-for-large-language-models","slug":"pac-prediction-sets-for-large-language-models","title":"PAC Prediction Sets for Large Language Models of Code","date":"2023-02-17","arxiv_id":"2302.08703","n_code_links":1,"syntology":null},{"paper":null,"slug":"foundation-models-for-natural-language","title":"Foundation Models for Natural Language Processing -- Pre-trained Language Models Integrating Media","date":"2023-02-16","arxiv_id":"2302.08575","n_code_links":0,"syntology":null},{"paper":null,"slug":"commonsense-reasoning-for-conversational-ai-a","title":"Commonsense Reasoning for Conversational AI: A Survey of the State of the Art","date":"2023-02-15","arxiv_id":"2302.07926","n_code_links":0,"syntology":null},{"paper":null,"slug":"artificial-intelligence-in-psychology","title":"Diminished Diversity-of-Thought in a Standard Large Language Model","date":"2023-02-13","arxiv_id":"2302.07267","n_code_links":0,"syntology":null},{"paper":null,"slug":"academic-writing-with-gpt-3-5-reflections-on","title":"Academic Writing with GPT-3.5: Reflections on Practices, Efficacy and Transparency","date":"2023-02-12","arxiv_id":"2304.11079","n_code_links":0,"syntology":null},{"paper":"/paper/combat-ai-with-ai-counteract-machine","slug":"combat-ai-with-ai-counteract-machine","title":"Combat AI With AI: Counteract Machine-Generated Fake Restaurant Reviews on Social Media","date":"2023-02-10","arxiv_id":"2302.07731","n_code_links":1,"syntology":null},{"paper":"/paper/the-wisdom-of-hindsight-makes-language-models","slug":"the-wisdom-of-hindsight-makes-language-models","title":"The Wisdom of Hindsight Makes Language Models Better Instruction Followers","date":"2023-02-10","arxiv_id":"2302.05206","n_code_links":1,"syntology":null},{"paper":"/paper/translating-natural-language-to-planning","slug":"translating-natural-language-to-planning","title":"Translating Natural Language to Planning Goals with Large-Language Models","date":"2023-02-10","arxiv_id":"2302.05128","n_code_links":1,"syntology":null},{"paper":null,"slug":"better-by-you-better-than-me-chatgpt3-as","title":"Better by you, better than me, chatgpt3 as writing assistance in students essays","date":"2023-02-09","arxiv_id":"2302.04536","n_code_links":0,"syntology":null},{"paper":null,"slug":"quantized-distributed-training-of-large","title":"Quantized Distributed Training of Large Models with Convergence Guarantees","date":"2023-02-05","arxiv_id":"2302.02390","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-language-reveals-about-perception","title":"Large language models predict human sensory judgments across six modalities","date":"2023-02-02","arxiv_id":"2302.01308","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-are-latent-variable-1","slug":"large-language-models-are-latent-variable-1","title":"Large Language Models Are Latent Variable Models: Explaining and Finding Good Demonstrations for In-Context Learning","date":"2023-01-27","arxiv_id":"2301.11916","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["wangxinyilinda/concept-based-demonstration-selection"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-stability-analysis-of-fine-tuning-a-pre","title":"A Stability Analysis of Fine-Tuning a Pre-Trained Model","date":"2023-01-24","arxiv_id":"2301.09820","n_code_links":0,"syntology":null},{"paper":"/paper/t2m-gpt-generating-human-motion-from-textual","slug":"t2m-gpt-generating-human-motion-from-textual","title":"T2M-GPT: Generating Human Motion from Textual Descriptions with Discrete Representations","date":"2023-01-15","arxiv_id":"2301.06052","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 4 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["Mael-zys/T2M-GPT"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":4,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/gpt-as-knowledge-worker-a-zero-shot","slug":"gpt-as-knowledge-worker-a-zero-shot","title":"GPT as Knowledge Worker: A Zero-Shot Evaluation of (AI)CPA Capabilities","date":"2023-01-11","arxiv_id":"2301.04408","n_code_links":1,"syntology":null},{"paper":null,"slug":"generating-human-motion-from-textual","title":"Generating Human Motion From Textual Descriptions With Discrete Representations","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/gpt-takes-the-bar-exam","slug":"gpt-takes-the-bar-exam","title":"GPT Takes the Bar Exam","date":"2022-12-29","arxiv_id":"2212.14402","n_code_links":5,"syntology":null},{"paper":"/paper/benchmark-for-uncertainty-robustness-in-self","slug":"benchmark-for-uncertainty-robustness-in-self","title":"Benchmark for Uncertainty & Robustness in Self-Supervised Learning","date":"2022-12-23","arxiv_id":"2212.12411","n_code_links":1,"syntology":null},{"paper":null,"slug":"kl-regularized-normalization-framework-for","title":"KL Regularized Normalization Framework for Low Resource Tasks","date":"2022-12-21","arxiv_id":"2212.11275","n_code_links":0,"syntology":null},{"paper":null,"slug":"true-detective-a-challenging-benchmark-for","title":"True Detective: A Deep Abductive Reasoning Benchmark Undoable for GPT-3 and Challenging for GPT-4","date":"2022-12-20","arxiv_id":"2212.10114","n_code_links":0,"syntology":null},{"paper":"/paper/why-can-gpt-learn-in-context-language-models","slug":"why-can-gpt-learn-in-context-language-models","title":"Why Can GPT Learn In-Context? Language Models Implicitly Perform Gradient Descent as Meta-Optimizers","date":"2022-12-20","arxiv_id":"2212.10559","n_code_links":1,"syntology":null},{"paper":"/paper/reasoning-with-language-model-prompting-a","slug":"reasoning-with-language-model-prompting-a","title":"Reasoning with Language Model Prompting: A Survey","date":"2022-12-19","arxiv_id":"2212.09597","n_code_links":2,"syntology":null},{"paper":null,"slug":"audio-driven-co-speech-gesture-video","title":"Audio-Driven Co-Speech Gesture Video Generation","date":"2022-12-05","arxiv_id":"2212.02350","n_code_links":0,"syntology":null},{"paper":null,"slug":"outfit-generation-and-recommendation-an","title":"Outfit Generation and Recommendation -- An Experimental Study","date":"2022-11-29","arxiv_id":"2211.16353","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-bloom-an-empirical-study-on","title":"Understanding BLOOM: An empirical study on diverse NLP tasks","date":"2022-11-27","arxiv_id":"2211.14865","n_code_links":0,"syntology":null},{"paper":"/paper/pointclip-v2-adapting-clip-for-powerful-3d","slug":"pointclip-v2-adapting-clip-for-powerful-3d","title":"PointCLIP V2: Prompting CLIP and GPT for Powerful 3D Open-world Learning","date":"2022-11-21","arxiv_id":"2211.11682","n_code_links":2,"syntology":{"ran":8,"of":12,"n_ran_checked":4,"n_instrument":4,"unverified":4,"pointer_only":4,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 1 violated, 3 with no contract checked; 4 where Syntology's instrument failed) · 4 unverified","official":{"repos":["yangyangyang127/pointclip_v2"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"conceptor-aided-debiasing-of-contextualized","title":"Conceptor-Aided Debiasing of Large Language Models","date":"2022-11-20","arxiv_id":"2211.11087","n_code_links":0,"syntology":null},{"paper":"/paper/random-ltd-random-and-layerwise-token","slug":"random-ltd-random-and-layerwise-token","title":"Random-LTD: Random and Layerwise Token Dropping Brings Efficient Training for Large-scale Transformers","date":"2022-11-17","arxiv_id":"2211.11586","n_code_links":1,"syntology":null},{"paper":"/paper/gptq-accurate-post-training-quantization-for","slug":"gptq-accurate-post-training-quantization-for","title":"GPTQ: Accurate Post-Training Quantization for Generative Pre-trained Transformers","date":"2022-10-31","arxiv_id":"2210.17323","n_code_links":17,"syntology":{"ran":5,"of":15,"n_ran_checked":2,"n_instrument":3,"unverified":10,"pointer_only":1,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 10 unverified","official":{"repos":["ist-daslab/gptq"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":3,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"trscore-a-novel-gpt-based-readability-scorer","title":"TRScore: A Novel GPT-based Readability Scorer for ASR Segmentation and Punctuation model evaluation and selection","date":"2022-10-27","arxiv_id":"2210.15104","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-robustness-of-prefix-tuning-in","title":"Exploring Robustness of Prefix Tuning in Noisy Data: A Case Study in Financial Sentiment Analysis","date":"2022-10-26","arxiv_id":"2211.05584","n_code_links":0,"syntology":null},{"paper":null,"slug":"ielm-an-open-information-extraction-benchmark","title":"IELM: An Open Information Extraction Benchmark for Pre-Trained Language Models","date":"2022-10-25","arxiv_id":"2210.14128","n_code_links":0,"syntology":null},{"paper":"/paper/emergent-world-representations-exploring-a","slug":"emergent-world-representations-exploring-a","title":"Emergent World Representations: Exploring a Sequence Model Trained on a Synthetic Task","date":"2022-10-24","arxiv_id":"2210.13382","n_code_links":4,"syntology":{"ran":6,"of":6,"n_ran_checked":5,"n_instrument":1,"unverified":0,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["likenneth/othello_world"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"meta-learning-pathologies-from-radiology","title":"Meta-learning Pathologies from Radiology Reports using Variance Aware Prototypical Networks","date":"2022-10-22","arxiv_id":"2210.13979","n_code_links":0,"syntology":null},{"paper":"/paper/a-causal-framework-to-quantify-the-robustness","slug":"a-causal-framework-to-quantify-the-robustness","title":"A Causal Framework to Quantify the Robustness of Mathematical Reasoning with Language Models","date":"2022-10-21","arxiv_id":"2210.12023","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":3,"n_instrument":2,"unverified":2,"pointer_only":7,"phrase":"5 ran (of which 2 constructed an object rather than computing a result; 3 with no instrument failure: 1 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["alestolfo/causal-math"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":2,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/general-image-descriptors-for-open-world","slug":"general-image-descriptors-for-open-world","title":"General Image Descriptors for Open World Image Retrieval using ViT CLIP","date":"2022-10-20","arxiv_id":"2210.11141","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ivanaer/g-universal-clip"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/biogpt-generative-pre-trained-transformer-for","slug":"biogpt-generative-pre-trained-transformer-for","title":"BioGPT: Generative Pre-trained Transformer for Biomedical Text Generation and Mining","date":"2022-10-19","arxiv_id":"2210.10341","n_code_links":4,"syntology":null},{"paper":null,"slug":"towards-a-neural-architecture-of-language","title":"Towards a neural architecture of language: Deep learning versus logistics of access in neural architectures for compositional processing","date":"2022-10-19","arxiv_id":"2210.10543","n_code_links":0,"syntology":null},{"paper":"/paper/dylora-parameter-efficient-tuning-of-pre","slug":"dylora-parameter-efficient-tuning-of-pre","title":"DyLoRA: Parameter Efficient Tuning of Pre-trained Models using Dynamic Search-Free Low-Rank Adaptation","date":"2022-10-14","arxiv_id":"2210.07558","n_code_links":2,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["huawei-noah/kd-nlp"],"state":"official: not harvested","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":[]}}},{"paper":"/paper/foundation-transformers","slug":"foundation-transformers","title":"Foundation Transformers","date":"2022-10-12","arxiv_id":"2210.06423","n_code_links":4,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["microsoft/unilm"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/exploiting-selection-bias-on-underspecified","slug":"exploiting-selection-bias-on-underspecified","title":"Underspecification in Language Modeling Tasks: A Causality-Informed Study of Gendered Pronoun Resolution","date":"2022-09-30","arxiv_id":"2210.00131","n_code_links":2,"syntology":null},{"paper":"/paper/using-large-language-models-to-simulate","slug":"using-large-language-models-to-simulate","title":"Using Large Language Models to Simulate Multiple Humans and Replicate Human Subject Studies","date":"2022-08-18","arxiv_id":"2208.10264","n_code_links":2,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 2 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified; every one of the 2 samples that ran constructed an object rather than computing a result","official":{"repos":["gatiaher/using-large-language-models-to-replicate-human-subject-studies"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/neural-embeddings-for-text","slug":"neural-embeddings-for-text","title":"Neural Embeddings for Text","date":"2022-08-17","arxiv_id":"2208.08386","n_code_links":1,"syntology":null},{"paper":"/paper/mocapact-a-multi-task-dataset-for-simulated","slug":"mocapact-a-multi-task-dataset-for-simulated","title":"MoCapAct: A Multi-Task Dataset for Simulated Humanoid Control","date":"2022-08-15","arxiv_id":"2208.07363","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["microsoft/MoCapAct"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"few-shot-training-llms-for-project-specific","title":"Few-shot training LLMs for project-specific code-summarization","date":"2022-07-09","arxiv_id":"2207.04237","n_code_links":0,"syntology":null},{"paper":null,"slug":"sensitivity-analysis-on-transferred-neural","title":"Sensitivity Analysis on Transferred Neural Architectures of BERT and GPT-2 for Financial Sentiment Analysis","date":"2022-07-07","arxiv_id":"2207.03037","n_code_links":0,"syntology":null},{"paper":null,"slug":"cocopie-xgen-a-full-stack-ai-oriented","title":"CoCoPIE XGen: A Full-Stack AI-Oriented Optimizing Framework","date":"2022-06-21","arxiv_id":"2206.10620","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-the-imitation-game-quantifying-and","slug":"beyond-the-imitation-game-quantifying-and","title":"Beyond the Imitation Game: Quantifying and extrapolating the capabilities of language models","date":"2022-06-09","arxiv_id":"2206.04615","n_code_links":6,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["google/BIG-bench"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/multi-agent-reinforcement-learning-is-a","slug":"multi-agent-reinforcement-learning-is-a","title":"Multi-Agent Reinforcement Learning is a Sequence Modeling Problem","date":"2022-05-30","arxiv_id":"2205.14953","n_code_links":1,"syntology":null},{"paper":"/paper/cped-a-large-scale-chinese-personalized-and-1","slug":"cped-a-large-scale-chinese-personalized-and-1","title":"CPED: A Large-Scale Chinese Personalized and Emotional Dialogue Dataset for Conversational AI","date":"2022-05-29","arxiv_id":"2205.14727","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-understanding-label-regularization","title":"Do we need Label Regularization to Fine-tune Pre-trained Language Models?","date":"2022-05-25","arxiv_id":"2205.12428","n_code_links":0,"syntology":null},{"paper":"/paper/transcormer-transformer-for-sentence-scoring","slug":"transcormer-transformer-for-sentence-scoring","title":"Transcormer: Transformer for Sentence Scoring with Sliding Language Modeling","date":"2022-05-25","arxiv_id":"2205.12986","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"on-the-role-of-bidirectionality-in-language","title":"On the Role of Bidirectionality in Language Model Pre-Training","date":"2022-05-24","arxiv_id":"2205.11726","n_code_links":0,"syntology":null},{"paper":"/paper/graphmae-self-supervised-masked-graph","slug":"graphmae-self-supervised-masked-graph","title":"GraphMAE: Self-Supervised Masked Graph Autoencoders","date":"2022-05-22","arxiv_id":"2205.10803","n_code_links":3,"syntology":{"ran":2,"of":3,"n_ran_checked":1,"n_instrument":1,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["thudm/graphmae"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/life-after-bert-what-do-other-muppets-2","slug":"life-after-bert-what-do-other-muppets-2","title":"Life after BERT: What do Other Muppets Understand about Language?","date":"2022-05-21","arxiv_id":"2205.10696","n_code_links":1,"syntology":null},{"paper":"/paper/prototypical-calibration-for-few-shot","slug":"prototypical-calibration-for-few-shot","title":"Prototypical Calibration for Few-shot Learning of Language Models","date":"2022-05-20","arxiv_id":"2205.10183","n_code_links":1,"syntology":null},{"paper":"/paper/automated-scoring-for-reading-comprehension","slug":"automated-scoring-for-reading-comprehension","title":"Automated Scoring for Reading Comprehension via In-context BERT Tuning","date":"2022-05-19","arxiv_id":"2205.09864","n_code_links":1,"syntology":null},{"paper":"/paper/provably-confidential-language-modelling-1","slug":"provably-confidential-language-modelling-1","title":"Provably Confidential Language Modelling","date":"2022-05-04","arxiv_id":"2205.01863","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["xuandongzhao/crt"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["found_in_text"]}}},{"paper":null,"slug":"exploring-artificial-intelligence-as-a","title":"Measuring artificial intelligence: a systematic assessment and implications for governance","date":"2022-04-21","arxiv_id":"2204.10304","n_code_links":0,"syntology":null},{"paper":null,"slug":"impact-of-tokenization-on-language-models-an-1","title":"Impact of Tokenization on Language Models: An Analysis for Turkish","date":"2022-04-19","arxiv_id":"2204.08832","n_code_links":0,"syntology":null},{"paper":"/paper/polling-latent-opinions-a-method-for-1","slug":"polling-latent-opinions-a-method-for-1","title":"Polling Latent Opinions: A Method for Computational Sociolinguistics Using Transformer Language Models","date":"2022-04-15","arxiv_id":"2204.07483","n_code_links":1,"syntology":null},{"paper":null,"slug":"foundationlayernorm-scaling-bert-and-gpt-to","title":"FoundationLayerNorm: Scaling BERT and GPT to 1,000 Layers","date":"2022-04-09","arxiv_id":"2204.04477","n_code_links":0,"syntology":null},{"paper":"/paper/data-augmentation-for-intent-classification-1","slug":"data-augmentation-for-intent-classification-1","title":"Data Augmentation for Intent Classification with Off-the-shelf Large Language Models","date":"2022-04-05","arxiv_id":"2204.01959","n_code_links":1,"syntology":{"ran":7,"of":10,"n_ran_checked":6,"n_instrument":1,"unverified":3,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":{"repos":["elementai/data-augmentation-with-llms"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"create-a-benchmark-for-chinese-short-video-1","title":"CREATE: A Benchmark for Chinese Short Video Retrieval and Title Generation","date":"2022-03-31","arxiv_id":"2203.16763","n_code_links":0,"syntology":null},{"paper":"/paper/bailando-3d-dance-generation-by-actor-critic","slug":"bailando-3d-dance-generation-by-actor-critic","title":"Bailando: 3D Dance Generation by Actor-Critic GPT with Choreographic Memory","date":"2022-03-24","arxiv_id":"2203.13055","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":4,"n_instrument":0,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 3 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["lisiyao21/bailando"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":3,"n_ran_no_instrument_failure":4,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"self-supervision-through-random-segments-with","title":"Self-supervision through Random Segments with Autoregressive Coding (RandSAC)","date":"2022-03-22","arxiv_id":"2203.12054","n_code_links":0,"syntology":null},{"paper":"/paper/grips-gradient-free-edit-based-instruction","slug":"grips-gradient-free-edit-based-instruction","title":"GrIPS: Gradient-free, Edit-based Instruction Search for Prompting Large Language Models","date":"2022-03-14","arxiv_id":"2203.07281","n_code_links":2,"syntology":{"ran":3,"of":4,"n_ran_checked":0,"n_instrument":3,"unverified":1,"pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["archiki/grips"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/elle-efficient-lifelong-pre-training-for-1","slug":"elle-efficient-lifelong-pre-training-for-1","title":"ELLE: Efficient Lifelong Pre-training for Emerging Data","date":"2022-03-12","arxiv_id":"2203.06311","n_code_links":1,"syntology":{"ran":4,"of":5,"n_ran_checked":1,"n_instrument":3,"unverified":1,"pointer_only":5,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["thunlp/elle"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}}],"record_sha256":"2d32d9a3b1f3020c0fa8c586b9d9017c200c912b8a70e3c20626622493395419","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}