{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/language-modelling/papers/131","list_of":"/task/language-modelling","task":"Language Modelling","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":131,"pages_in_order":177,"rows_per_page":100,"rows":[13001,13100],"of":17610,"counts":{"archive_papers_tagged":17610,"with_a_code_link":7012,"where_syntology_ran_a_sample":2428,"not_listed_spam_title":0,"listed":17610,"listed_where_code_ran":2428,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2027,"every_run_a_failure_of_syntologys_instrument":401,"listed_with_a_run_with_no_instrument_failure":2027,"listed_every_run_a_failure_of_syntologys_instrument":401,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/language-modelling","prev":"/task/language-modelling/papers/130","next":"/task/language-modelling/papers/132","papers":[{"url":null,"slug":"unified-text-structuralization-with","title":"Unified Text Structuralization with Instruction-tuned Language Models","date":"2023-03-27","arxiv_id":"2303.14956","repositories_listed":0,"syntology":null},{"url":null,"slug":"backdoor-attacks-with-input-unique-triggers","title":"Backdoor Attacks with Input-unique Triggers in NLP","date":"2023-03-25","arxiv_id":"2303.14325","repositories_listed":0,"syntology":null},{"url":null,"slug":"sem4sap-synonymous-expression-mining-from","title":"Sem4SAP: Synonymous Expression Mining From Open Knowledge Graph For Language Model Synonym-Aware Pretraining","date":"2023-03-25","arxiv_id":"2303.14425","repositories_listed":0,"syntology":null},{"url":null,"slug":"toward-open-domain-slot-filling-via-self","title":"Toward Open-domain Slot Filling via Self-supervised Co-training","date":"2023-03-24","arxiv_id":"2303.13801","repositories_listed":0,"syntology":null},{"url":null,"slug":"unleasing-chatgpt-on-the-metaverse-savior-or","title":"Unleashing GPT on the Metaverse: Savior or Destroyer?","date":"2023-03-24","arxiv_id":"2303.13856","repositories_listed":0,"syntology":null},{"url":null,"slug":"attention-based-speech-enhancement-using","title":"Attention-based Speech Enhancement Using Human Quality Perception Modelling","date":"2023-03-23","arxiv_id":"2303.13685","repositories_listed":0,"syntology":null},{"url":null,"slug":"chatgpt-for-shaping-the-future-of-dentistry","title":"ChatGPT for Shaping the Future of Dentistry: The Potential of Multi-Modal Large Language Model","date":"2023-03-23","arxiv_id":"2304.03086","repositories_listed":0,"syntology":null},{"url":null,"slug":"spec-a-soft-prompt-based-calibration-on","title":"SPeC: A Soft Prompt-Based Calibration on Performance Variability of Large Language Model in Clinical Notes Summarization","date":"2023-03-23","arxiv_id":"2303.13035","repositories_listed":0,"syntology":null},{"url":null,"slug":"three-ways-to-improve-feature-alignment-for","title":"Three ways to improve feature alignment for open vocabulary detection","date":"2023-03-23","arxiv_id":"2303.13518","repositories_listed":0,"syntology":null},{"url":null,"slug":"can-we-trust-the-evaluation-on-chatgpt","title":"Can we trust the evaluation on ChatGPT?","date":"2023-03-22","arxiv_id":"2303.12767","repositories_listed":0,"syntology":null},{"url":null,"slug":"extracting-physical-rehabilitation-exercise","title":"Mining Clinical Notes for Physical Rehabilitation Exercise Information: Natural Language Processing Algorithm Development and Validation Study","date":"2023-03-22","arxiv_id":"2303.13466","repositories_listed":0,"syntology":null},{"url":null,"slug":"frozen-language-model-helps-ecg-zero-shot","title":"Frozen Language Model Helps ECG Zero-Shot Learning","date":"2023-03-22","arxiv_id":"2303.12311","repositories_listed":0,"syntology":null},{"url":null,"slug":"salient-span-masking-for-temporal","title":"Salient Span Masking for Temporal Understanding","date":"2023-03-22","arxiv_id":"2303.12860","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-complete-survey-on-generative-ai-aigc-is","title":"A Complete Survey on Generative AI (AIGC): Is ChatGPT from GPT-4 to GPT-5 All You Need?","date":"2023-03-21","arxiv_id":"2303.11717","repositories_listed":0,"syntology":null},{"url":null,"slug":"hop-history-enhanced-and-order-aware-pre","title":"HOP+: History-enhanced and Order-aware Pre-training for Vision-and-Language Navigation","date":"2023-03-20","arxiv_id":null,"repositories_listed":0,"syntology":null},{"url":null,"slug":"large-language-models-and-simple-stupid-bugs","title":"Large Language Models and Simple, Stupid Bugs","date":"2023-03-20","arxiv_id":"2303.11455","repositories_listed":0,"syntology":null},{"url":null,"slug":"maximizing-penetration-testing-success-with","title":"Maximizing Penetration Testing Success with Effective Reconnaissance Techniques using ChatGPT","date":"2023-03-20","arxiv_id":"2307.06391","repositories_listed":0,"syntology":null},{"url":null,"slug":"mind-meets-machine-unravelling-gpt-4-s","title":"Mind meets machine: Unravelling GPT-4's cognitive psychology","date":"2023-03-20","arxiv_id":"2303.11436","repositories_listed":0,"syntology":null},{"url":null,"slug":"multimodal-shannon-game-with-images","title":"Multimodal Shannon Game with Images","date":"2023-03-20","arxiv_id":"2303.11192","repositories_listed":0,"syntology":null},{"url":null,"slug":"on-the-fly-text-retrieval-for-end-to-end-asr","title":"On-the-fly Text Retrieval for End-to-End ASR Adaptation","date":"2023-03-20","arxiv_id":"2303.10942","repositories_listed":0,"syntology":null},{"url":null,"slug":"pangu-s-towards-trillion-parameter-language","title":"PanGu-Σ: Towards Trillion Parameter Language Model with Sparse Heterogeneous Computing","date":"2023-03-20","arxiv_id":"2303.10845","repositories_listed":0,"syntology":null},{"url":null,"slug":"label-name-is-mantra-unifying-point-cloud","title":"Label Name is Mantra: Unifying Point Cloud Segmentation across Heterogeneous Datasets","date":"2023-03-19","arxiv_id":"2303.10585","repositories_listed":0,"syntology":null},{"url":null,"slug":"lightweight-contrastive-protein-structure","title":"CCPL: Cross-modal Contrastive Protein Learning","date":"2023-03-19","arxiv_id":"2303.11783","repositories_listed":0,"syntology":null},{"url":null,"slug":"doric-domain-robust-fine-tuning-for-open","title":"DORIC : Domain Robust Fine-Tuning for Open Intent Clustering through Dependency Parsing","date":"2023-03-17","arxiv_id":"2303.09827","repositories_listed":0,"syntology":null},{"url":"/paper/smartbert-a-promotion-of-dynamic-early","slug":"smartbert-a-promotion-of-dynamic-early","title":"SmartBERT: A Promotion of Dynamic Early Exiting Mechanism for Accelerating BERT Inference","date":"2023-03-16","arxiv_id":"2303.09266","repositories_listed":0,"syntology":{"n":4,"n_ran":3,"n_constructed":3,"n_ran_checked":3,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":3,"n_pointer_only":0,"phrase":"3 ran (of which 3 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified; every one of the 3 samples that ran constructed an object rather than computing a result","sample_list":"/paper/smartbert-a-promotion-of-dynamic-early#ran","syntology_url":"https://syntology.ai/paper/2303.09266","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.09266"}},"official":null}},{"url":null,"slug":"towards-the-scalable-evaluation-of","title":"Towards the Scalable Evaluation of Cooperativeness in Language Models","date":"2023-03-16","arxiv_id":"2303.13360","repositories_listed":0,"syntology":null},{"url":null,"slug":"translating-radiology-reports-into-plain","title":"Translating Radiology Reports into Plain Language using ChatGPT and GPT-4 with Prompt Learning: Promising Results, Limitations, and Potential","date":"2023-03-16","arxiv_id":"2303.09038","repositories_listed":0,"syntology":null},{"url":null,"slug":"chatgpt-or-grammarly-evaluating-chatgpt-on","title":"ChatGPT or Grammarly? Evaluating ChatGPT on Grammatical Error Correction Benchmark","date":"2023-03-15","arxiv_id":"2303.13648","repositories_listed":0,"syntology":null},{"url":null,"slug":"contextualized-medication-information","title":"Contextualized Medication Information Extraction Using Transformer-based Deep Learning Architectures","date":"2023-03-14","arxiv_id":"2303.08259","repositories_listed":0,"syntology":null},{"url":null,"slug":"do-transformers-parse-while-predicting-the","title":"Do Transformers Parse while Predicting the Masked Word?","date":"2023-03-14","arxiv_id":"2303.08117","repositories_listed":0,"syntology":null},{"url":null,"slug":"finding-the-needle-in-a-haystack-unsupervised","title":"Finding the Needle in a Haystack: Unsupervised Rationale Extraction from Long Text Classifiers","date":"2023-03-14","arxiv_id":"2303.07991","repositories_listed":0,"syntology":null},{"url":null,"slug":"simfluence-modeling-the-influence-of","title":"Simfluence: Modeling the Influence of Individual Training Examples by Simulating Training Runs","date":"2023-03-14","arxiv_id":"2303.08114","repositories_listed":0,"syntology":null},{"url":null,"slug":"generating-multiple-choice-questions-for","title":"Generating multiple-choice questions for medical question answering with distractors and cue-masking","date":"2023-03-13","arxiv_id":"2303.07069","repositories_listed":0,"syntology":null},{"url":null,"slug":"odin-on-demand-data-formulation-to-mitigate","title":"ODIN: On-demand Data Formulation to Mitigate Dataset Lock-in","date":"2023-03-13","arxiv_id":"2303.06832","repositories_listed":0,"syntology":null},{"url":null,"slug":"consistency-analysis-of-chatgpt","title":"Consistency Analysis of ChatGPT","date":"2023-03-11","arxiv_id":"2303.06273","repositories_listed":0,"syntology":null},{"url":null,"slug":"learning-combinatorial-prompts-for-universal","title":"Learning Combinatorial Prompts for Universal Controllable Image Captioning","date":"2023-03-11","arxiv_id":"2303.06338","repositories_listed":0,"syntology":null},{"url":null,"slug":"algorithmic-ghost-in-the-research-shell-large","title":"Algorithmic Ghost in the Research Shell: Large Language Models and Academic Knowledge Creation in Management Research","date":"2023-03-10","arxiv_id":"2303.07304","repositories_listed":0,"syntology":null},{"url":null,"slug":"an-overview-on-language-models-recent","title":"An Overview on Language Models: Recent Developments and Outlook","date":"2023-03-10","arxiv_id":"2303.05759","repositories_listed":0,"syntology":null},{"url":null,"slug":"rewarding-chatbots-for-real-world-engagement","title":"Rewarding Chatbots for Real-World Engagement with Millions of Users","date":"2023-03-10","arxiv_id":"2303.06135","repositories_listed":0,"syntology":null},{"url":null,"slug":"susceptibility-to-influence-of-large-language","title":"Susceptibility to Influence of Large Language Models","date":"2023-03-10","arxiv_id":"2303.06074","repositories_listed":0,"syntology":null},{"url":null,"slug":"towards-moe-deployment-mitigating","title":"Towards MoE Deployment: Mitigating Inefficiencies in Mixture-of-Expert (MoE) Inference","date":"2023-03-10","arxiv_id":"2303.06182","repositories_listed":0,"syntology":null},{"url":"/paper/can-a-frozen-pretrained-language-model-be","slug":"can-a-frozen-pretrained-language-model-be","title":"Can a Frozen Pretrained Language Model be used for Zero-shot Neural Retrieval on Entity-centric Questions?","date":"2023-03-09","arxiv_id":"2303.05153","repositories_listed":0,"syntology":null},{"url":null,"slug":"knowledge-augmented-few-shot-visual-relation","title":"Knowledge-augmented Few-shot Visual Relation Detection","date":"2023-03-09","arxiv_id":"2303.05342","repositories_listed":0,"syntology":null},{"url":null,"slug":"planning-with-large-language-models-for-code","title":"Planning with Large Language Models for Code Generation","date":"2023-03-09","arxiv_id":"2303.05510","repositories_listed":0,"syntology":null},{"url":null,"slug":"replacement-as-a-self-supervision-for-fine","title":"Refined Vision-Language Modeling for Fine-grained Multi-modal Pre-training","date":"2023-03-09","arxiv_id":"2303.05313","repositories_listed":0,"syntology":null},{"url":null,"slug":"weakly-supervised-hoi-detection-from","title":"Weakly-Supervised HOI Detection from Interaction Labels Only and Language/Vision-Language Priors","date":"2023-03-09","arxiv_id":"2303.05546","repositories_listed":0,"syntology":null},{"url":null,"slug":"extending-the-pre-training-of-bloom-for","title":"Extending the Pre-Training of BLOOM for Improved Support of Traditional Chinese: Models, Methods and Results","date":"2023-03-08","arxiv_id":"2303.04715","repositories_listed":0,"syntology":null},{"url":null,"slug":"magnushammer-a-transformer-based-approach-to","title":"Magnushammer: A Transformer-Based Approach to Premise Selection","date":"2023-03-08","arxiv_id":"2303.04488","repositories_listed":0,"syntology":null},{"url":null,"slug":"chatgpt-beginning-of-an-end-of-manual","title":"ChatGPT: Beginning of an End of Manual Linguistic Data Annotation? Use Case of Automatic Genre Identification","date":"2023-03-07","arxiv_id":"2303.03953","repositories_listed":0,"syntology":null},{"url":null,"slug":"german-bert-model-for-legal-named-entity","title":"German BERT Model for Legal Named Entity Recognition","date":"2023-03-07","arxiv_id":"2303.05388","repositories_listed":0,"syntology":null},{"url":null,"slug":"making-a-computational-attorney","title":"Making a Computational Attorney","date":"2023-03-07","arxiv_id":"2303.05383","repositories_listed":0,"syntology":null},{"url":"/paper/the-bigscience-roots-corpus-a-1-6tb-composite","slug":"the-bigscience-roots-corpus-a-1-6tb-composite","title":"The BigScience ROOTS Corpus: A 1.6TB Composite Multilingual Dataset","date":"2023-03-07","arxiv_id":"2303.03915","repositories_listed":0,"syntology":{"n":4,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/the-bigscience-roots-corpus-a-1-6tb-composite#ran","syntology_url":"https://syntology.ai/paper/2303.03915","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2303.03915"}},"official":null}},{"url":null,"slug":"chatgpt-is-on-the-horizon-could-a-large","title":"ChatGPT is on the Horizon: Could a Large Language Model be Suitable for Intelligent Traffic Safety Research and Applications?","date":"2023-03-06","arxiv_id":"2303.05382","repositories_listed":0,"syntology":null},{"url":null,"slug":"data-portraits-recording-foundation-model","title":"Data Portraits: Recording Foundation Model Training Data","date":"2023-03-06","arxiv_id":"2303.03919","repositories_listed":0,"syntology":null},{"url":null,"slug":"foundationtts-text-to-speech-for-asr","title":"FoundationTTS: Text-to-Speech for ASR Customization with Generative Language Model","date":"2023-03-06","arxiv_id":"2303.02939","repositories_listed":0,"syntology":null},{"url":null,"slug":"model-agnostic-meta-learning-for-natural","title":"Model-Agnostic Meta-Learning for Natural Language Understanding Tasks in Finance","date":"2023-03-06","arxiv_id":"2303.02841","repositories_listed":0,"syntology":null},{"url":null,"slug":"spelling-convention-sensitivity-in-neural","title":"Spelling convention sensitivity in neural language models","date":"2023-03-06","arxiv_id":"2303.03457","repositories_listed":0,"syntology":null},{"url":null,"slug":"could-a-large-language-model-be-conscious","title":"Could a Large Language Model be Conscious?","date":"2023-03-04","arxiv_id":"2303.07103","repositories_listed":0,"syntology":null},{"url":null,"slug":"end-to-end-speech-recognition-a-survey","title":"End-to-End Speech Recognition: A Survey","date":"2023-03-03","arxiv_id":"2303.03329","repositories_listed":0,"syntology":null},{"url":null,"slug":"exploiting-language-relatedness-in-machine","title":"Exploiting Language Relatedness in Machine Translation Through Domain Adaptation Techniques","date":"2023-03-03","arxiv_id":"2303.01793","repositories_listed":0,"syntology":null},{"url":null,"slug":"reprem-representation-pre-training-with","title":"RePreM: Representation Pre-training with Masked Model for Reinforcement Learning","date":"2023-03-03","arxiv_id":"2303.01668","repositories_listed":0,"syntology":null},{"url":null,"slug":"will-affective-computing-emerge-from","title":"Will Affective Computing Emerge from Foundation Models and General AI? A First Evaluation on ChatGPT","date":"2023-03-03","arxiv_id":"2303.03186","repositories_listed":0,"syntology":null},{"url":null,"slug":"ai-and-the-fci-can-chatgpt-project-an","title":"AI and the FCI: Can ChatGPT Project an Understanding of Introductory Physics?","date":"2023-03-02","arxiv_id":"2303.01067","repositories_listed":0,"syntology":null},{"url":null,"slug":"benchdirect-a-directed-language-model-for","title":"BenchDirect: A Directed Language Model for Compiler Benchmarks","date":"2023-03-02","arxiv_id":"2303.01557","repositories_listed":0,"syntology":null},{"url":null,"slug":"how-will-language-modelers-like-chatgpt","title":"How will Language Modelers like ChatGPT Affect Occupations and Industries?","date":"2023-03-02","arxiv_id":"2303.01157","repositories_listed":0,"syntology":null},{"url":null,"slug":"open-world-object-manipulation-using-pre","title":"Open-World Object Manipulation using Pre-trained Vision-Language Models","date":"2023-03-02","arxiv_id":"2303.00905","repositories_listed":0,"syntology":null},{"url":null,"slug":"semiparametric-language-models-are-scalable","title":"Semiparametric Language Models Are Scalable Continual Learners","date":"2023-03-02","arxiv_id":"2303.01421","repositories_listed":0,"syntology":null},{"url":null,"slug":"almanac-knowledge-grounded-language-models","title":"Almanac: Retrieval-Augmented Language Models for Clinical Medicine","date":"2023-03-01","arxiv_id":"2303.01229","repositories_listed":0,"syntology":null},{"url":null,"slug":"domain-adapted-large-language-models-for","title":"Domain-adapted large language models for classifying nuclear medicine reports","date":"2023-03-01","arxiv_id":"2303.01258","repositories_listed":0,"syntology":null},{"url":null,"slug":"grounded-decoding-guiding-text-generation","title":"Grounded Decoding: Guiding Text Generation with Grounded Models for Embodied Agents","date":"2023-03-01","arxiv_id":"2303.00855","repositories_listed":0,"syntology":null},{"url":null,"slug":"n-best-t5-robust-asr-error-correction-using","title":"N-best T5: Robust ASR Error Correction using Multiple Input Hypotheses and Constrained Decoding Space","date":"2023-03-01","arxiv_id":"2303.00456","repositories_listed":0,"syntology":null},{"url":"/paper/speechprompt-v2-prompt-tuning-for-speech","slug":"speechprompt-v2-prompt-tuning-for-speech","title":"SpeechPrompt v2: Prompt Tuning for Speech Classification Tasks","date":"2023-03-01","arxiv_id":"2303.00733","repositories_listed":0,"syntology":null},{"url":null,"slug":"efficient-masked-autoencoders-with-self","title":"Efficient Masked Autoencoders with Self-Consistency","date":"2023-02-28","arxiv_id":"2302.14431","repositories_listed":0,"syntology":null},{"url":null,"slug":"weighted-sampling-for-masked-language","title":"Weighted Sampling for Masked Language Modeling","date":"2023-02-28","arxiv_id":"2302.14225","repositories_listed":0,"syntology":null},{"url":null,"slug":"duration-aware-pause-insertion-using-pre","title":"Duration-aware pause insertion using pre-trained language model for multi-speaker text-to-speech","date":"2023-02-27","arxiv_id":"2302.13652","repositories_listed":0,"syntology":null},{"url":null,"slug":"finding-supporting-examples-for-in-context","title":"Finding Support Examples for In-Context Learning","date":"2023-02-27","arxiv_id":"2302.13539","repositories_listed":0,"syntology":null},{"url":null,"slug":"leveraging-large-language-model-and-story","title":"Leveraging Large Language Model and Story-Based Gamification in Intelligent Tutoring System to Scaffold Introductory Programming Courses: A Design-Based Research Study","date":"2023-02-25","arxiv_id":"2302.12834","repositories_listed":0,"syntology":null},{"url":null,"slug":"topic-selective-graph-network-for-topic","title":"Topic-Selective Graph Network for Topic-Focused Summarization","date":"2023-02-25","arxiv_id":"2302.13106","repositories_listed":0,"syntology":null},{"url":null,"slug":"toward-fairness-in-text-generation-via-mutual","title":"Toward Fairness in Text Generation via Mutual Information Minimization based on Importance Sampling","date":"2023-02-25","arxiv_id":"2302.13136","repositories_listed":0,"syntology":null},{"url":null,"slug":"factual-consistency-oriented-speech","title":"Factual Consistency Oriented Speech Recognition","date":"2023-02-24","arxiv_id":"2302.12369","repositories_listed":0,"syntology":null},{"url":null,"slug":"generative-sentiment-transfer-via-adaptive","title":"Generative Sentiment Transfer via Adaptive Masking","date":"2023-02-23","arxiv_id":"2302.12045","repositories_listed":0,"syntology":null},{"url":null,"slug":"on-the-generalization-ability-of-retrieval","title":"On the Generalization Ability of Retrieval-Enhanced Transformers","date":"2023-02-23","arxiv_id":"2302.12128","repositories_listed":0,"syntology":null},{"url":"/paper/vlsp-2022-evjvqa-challenge-multilingual","slug":"vlsp-2022-evjvqa-challenge-multilingual","title":"EVJVQA Challenge: Multilingual Visual Question Answering","date":"2023-02-23","arxiv_id":"2302.11752","repositories_listed":0,"syntology":null},{"url":null,"slug":"badgpt-exploring-security-vulnerabilities-of","title":"BadGPT: Exploring Security Vulnerabilities of ChatGPT via Backdoor Attacks to InstructGPT","date":"2023-02-21","arxiv_id":"2304.12298","repositories_listed":0,"syntology":null},{"url":null,"slug":"k-nn-adapter-efficient-domain-adaptation-for","title":"$k$NN-Adapter: Efficient Domain Adaptation for Black-Box Language Models","date":"2023-02-21","arxiv_id":"2302.10879","repositories_listed":0,"syntology":null},{"url":null,"slug":"playing-the-werewolf-game-with-artificial","title":"Playing the Werewolf game with artificial intelligence for language understanding","date":"2023-02-21","arxiv_id":"2302.10646","repositories_listed":0,"syntology":null},{"url":null,"slug":"emphasizing-unseen-words-new-vocabulary","title":"Emphasizing Unseen Words: New Vocabulary Acquisition for End-to-End Speech Recognition","date":"2023-02-20","arxiv_id":"2302.09723","repositories_listed":0,"syntology":null},{"url":null,"slug":"stoa-vlp-spatial-temporal-modeling-of-object","title":"STOA-VLP: Spatial-Temporal Modeling of Object and Action for Video-Language Pre-training","date":"2023-02-20","arxiv_id":"2302.09736","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-comprehensive-survey-on-pretrained","title":"A Comprehensive Survey on Pretrained Foundation Models: A History from BERT to ChatGPT","date":"2023-02-18","arxiv_id":"2302.09419","repositories_listed":0,"syntology":null},{"url":null,"slug":"bag-of-tricks-for-effective-language-model","title":"Bag of Tricks for Effective Language Model Pretraining and Downstream Adaptation: A Case Study on GLUE","date":"2023-02-18","arxiv_id":"2302.09268","repositories_listed":0,"syntology":null},{"url":null,"slug":"a-simplistic-model-of-neural-scaling-laws","title":"Multiperiodic Processes: Ergodic Sources with a Sublinear Entropy","date":"2023-02-17","arxiv_id":"2302.09049","repositories_listed":0,"syntology":null},{"url":null,"slug":"entry-separation-using-a-mixed-visual-and","title":"Entry Separation using a Mixed Visual and Textual Language Model: Application to 19th century French Trade Directories","date":"2023-02-17","arxiv_id":"2302.08948","repositories_listed":0,"syntology":null},{"url":null,"slug":"gpt4mia-utilizing-geneative-pre-trained","title":"GPT4MIA: Utilizing Generative Pre-trained Transformer (GPT-3) as A Plug-and-Play Transductive Model for Medical Image Analysis","date":"2023-02-17","arxiv_id":"2302.08722","repositories_listed":0,"syntology":null},{"url":null,"slug":"massively-multilingual-shallow-fusion-with","title":"Massively Multilingual Shallow Fusion with Large Language Models","date":"2023-02-17","arxiv_id":"2302.08917","repositories_listed":0,"syntology":null},{"url":null,"slug":"privately-customizing-prefinetuning-to-better","title":"Privately Customizing Prefinetuning to Better Match User Data in Federated Learning","date":"2023-02-17","arxiv_id":"2302.09042","repositories_listed":0,"syntology":null},{"url":null,"slug":"prompting-large-language-models-with-the","title":"Prompting Large Language Models With the Socratic Method","date":"2023-02-17","arxiv_id":"2303.08769","repositories_listed":0,"syntology":null},{"url":null,"slug":"adaptable-end-to-end-asr-models-using","title":"Adaptable End-to-End ASR Models using Replaceable Internal LMs and Residual Softmax","date":"2023-02-16","arxiv_id":"2302.08579","repositories_listed":0,"syntology":null},{"url":null,"slug":"bridge-the-gap-between-language-models-and","title":"Bridge the Gap between Language models and Tabular Understanding","date":"2023-02-16","arxiv_id":"2302.09302","repositories_listed":0,"syntology":null},{"url":null,"slug":"jeit-joint-end-to-end-model-and-internal","title":"JEIT: Joint End-to-End Model and Internal Language Model Training for Speech Recognition","date":"2023-02-16","arxiv_id":"2302.08583","repositories_listed":0,"syntology":null},{"url":null,"slug":"labelprompt-effective-prompt-based-learning","title":"LabelPrompt: Effective Prompt-based Learning for Relation Classification","date":"2023-02-16","arxiv_id":"2302.08068","repositories_listed":0,"syntology":null}],"record_sha256":"2ab4dbd3062b5d1393a6a7c4ef818367592a0cc600cc8d98bcd946233e64dd87","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}