{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/sentencepiece/papers/3","list_of":"/method/sentencepiece","method":"SentencePiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":3,"pages_in_order":10,"rows_per_page":100,"rows":[201,300],"of":908,"counts":{"archive_papers_tagged":908,"with_a_code_link":438,"where_syntology_ran_a_sample":111,"not_listed_spam_title":0,"listed":908,"listed_where_code_ran":111,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":96,"every_run_a_failure_of_syntologys_instrument":15,"listed_with_a_run_with_no_instrument_failure":96,"listed_every_run_a_failure_of_syntologys_instrument":15,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/sentencepiece","prev":"/method/sentencepiece/papers/2","next":"/method/sentencepiece/papers/4","papers":[{"paper":null,"slug":"a-synthetic-data-approach-for-domain","title":"A synthetic data approach for domain generalization of NLI models","date":"2024-02-19","arxiv_id":"2402.12368","n_code_links":0,"syntology":null},{"paper":null,"slug":"key-ingredients-for-effective-zero-shot-cross","title":"Key ingredients for effective zero-shot cross-lingual knowledge transfer in generative tasks","date":"2024-02-19","arxiv_id":"2402.12279","n_code_links":0,"syntology":null},{"paper":null,"slug":"emerging-opportunities-of-using-large","title":"Emerging Opportunities of Using Large Language Models for Translation Between Drug Molecules and Indications","date":"2024-02-14","arxiv_id":"2402.09588","n_code_links":0,"syntology":null},{"paper":null,"slug":"fgeo-tp-a-language-model-enhanced-solver-for","title":"FGeo-TP: A Language Model-Enhanced Solver for Geometry Problems","date":"2024-02-14","arxiv_id":"2402.09047","n_code_links":0,"syntology":null},{"paper":null,"slug":"eliciting-big-five-personality-traits-in","title":"Eliciting Personality Traits in Large Language Models","date":"2024-02-13","arxiv_id":"2402.08341","n_code_links":0,"syntology":null},{"paper":"/paper/improving-black-box-robustness-with-in","slug":"improving-black-box-robustness-with-in","title":"Improving Black-box Robustness with In-Context Rewriting","date":"2024-02-13","arxiv_id":"2402.08225","n_code_links":1,"syntology":null},{"paper":"/paper/inksight-offline-to-online-handwriting","slug":"inksight-offline-to-online-handwriting","title":"InkSight: Offline-to-Online Handwriting Conversion by Learning to Read and Write","date":"2024-02-08","arxiv_id":"2402.05804","n_code_links":1,"syntology":null},{"paper":null,"slug":"lens-a-foundation-model-for-network-traffic","title":"Lens: A Foundation Model for Network Traffic","date":"2024-02-06","arxiv_id":"2402.03646","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-a-unified-language-model-for","title":"CorpusLM: Towards a Unified Language Model on Corpus for Knowledge-Intensive Tasks","date":"2024-02-02","arxiv_id":"2402.01176","n_code_links":0,"syntology":null},{"paper":null,"slug":"what-will-my-model-forget-forecasting","title":"What Will My Model Forget? Forecasting Forgotten Examples in Language Model Refinement","date":"2024-02-02","arxiv_id":"2402.01865","n_code_links":0,"syntology":null},{"paper":"/paper/improving-semantic-control-in-discrete-latent","slug":"improving-semantic-control-in-discrete-latent","title":"Improving Semantic Control in Discrete Latent Spaces with Transformer Quantized Variational Autoencoders","date":"2024-02-01","arxiv_id":"2402.00723","n_code_links":1,"syntology":null},{"paper":"/paper/breaking-free-transformer-models-task","slug":"breaking-free-transformer-models-task","title":"Breaking Free Transformer Models: Task-specific Context Attribution Promises Improved Generalizability Without Fine-tuning Pre-trained LLMs","date":"2024-01-30","arxiv_id":"2401.16638","n_code_links":1,"syntology":null},{"paper":"/paper/topro-token-level-prompt-decomposition-for","slug":"topro-token-level-prompt-decomposition-for","title":"ToPro: Token-Level Prompt Decomposition for Cross-Lingual Sequence Labeling Tasks","date":"2024-01-29","arxiv_id":"2401.16589","n_code_links":1,"syntology":null},{"paper":null,"slug":"scalable-link-prediction-on-large-scale","title":"LPNL: Scalable Link Prediction with Large Language Models","date":"2024-01-24","arxiv_id":"2401.13227","n_code_links":0,"syntology":null},{"paper":"/paper/apt-adaptive-pruning-and-tuning-pretrained","slug":"apt-adaptive-pruning-and-tuning-pretrained","title":"APT: Adaptive Pruning and Tuning Pretrained Language Models for Efficient Training and Inference","date":"2024-01-22","arxiv_id":"2401.12200","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":5,"n_instrument":1,"unverified":1,"pointer_only":1,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["roim1998/apt"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"the-right-model-for-the-job-an-evaluation-of","title":"The Right Model for the Job: An Evaluation of Legal Multi-Label Classification Baselines","date":"2024-01-22","arxiv_id":"2401.11852","n_code_links":0,"syntology":null},{"paper":"/paper/finding-a-needle-in-the-adversarial-haystack","slug":"finding-a-needle-in-the-adversarial-haystack","title":"Finding a Needle in the Adversarial Haystack: A Targeted Paraphrasing Approach For Uncovering Edge Cases with Minimal Distribution Distortion","date":"2024-01-21","arxiv_id":"2401.11373","n_code_links":1,"syntology":null},{"paper":"/paper/langbridge-multilingual-reasoning-without","slug":"langbridge-multilingual-reasoning-without","title":"LangBridge: Multilingual Reasoning Without Multilingual Supervision","date":"2024-01-19","arxiv_id":"2401.10695","n_code_links":1,"syntology":{"ran":5,"of":12,"n_ran_checked":3,"n_instrument":2,"unverified":7,"pointer_only":12,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 2 where Syntology's instrument failed) · 7 unverified","official":{"repos":["kaistAI/LangBridge"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":7,"ran_from_kinds":["official"]}}},{"paper":"/paper/deciphering-textual-authenticity-a","slug":"deciphering-textual-authenticity-a","title":"Deciphering Textual Authenticity: A Generalized Strategy through the Lens of Large Language Semantics for Detecting Human vs. Machine-Generated Text","date":"2024-01-17","arxiv_id":"2401.09407","n_code_links":1,"syntology":null},{"paper":"/paper/mapping-transformer-leveraged-embeddings-for","slug":"mapping-transformer-leveraged-embeddings-for","title":"Mapping Transformer Leveraged Embeddings for Cross-Lingual Document Representation","date":"2024-01-12","arxiv_id":"2401.06583","n_code_links":1,"syntology":null},{"paper":"/paper/pizzacommonsense-learning-to-model","slug":"pizzacommonsense-learning-to-model","title":"PizzaCommonSense: Learning to Model Commonsense Reasoning about Intermediate Steps in Cooking Recipes","date":"2024-01-12","arxiv_id":"2401.06930","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-assessment-on-comprehending-mental-health","title":"An Assessment on Comprehending Mental Health through Large Language Models","date":"2024-01-09","arxiv_id":"2401.04592","n_code_links":0,"syntology":null},{"paper":"/paper/ast-t5-structure-aware-pretraining-for-code","slug":"ast-t5-structure-aware-pretraining-for-code","title":"AST-T5: Structure-Aware Pretraining for Code Generation and Understanding","date":"2024-01-05","arxiv_id":"2401.03003","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["gonglinyuan/ast_t5"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"are-llms-robust-for-spoken-dialogues","title":"Are LLMs Robust for Spoken Dialogues?","date":"2024-01-04","arxiv_id":"2401.02297","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-in-mental-health-care-a","title":"Large Language Models in Mental Health Care: a Scoping Review","date":"2024-01-01","arxiv_id":"2401.02984","n_code_links":0,"syntology":null},{"paper":"/paper/scaling-down-litting-up-efficient-zero-shot","slug":"scaling-down-litting-up-efficient-zero-shot","title":"Scaling Down, LiTting Up: Efficient Zero-Shot Listwise Reranking with Seq2seq Encoder-Decoder Models","date":"2023-12-26","arxiv_id":"2312.16098","n_code_links":2,"syntology":null},{"paper":"/paper/numerical-reasoning-for-financial-reports","slug":"numerical-reasoning-for-financial-reports","title":"Numerical Reasoning for Financial Reports","date":"2023-12-22","arxiv_id":"2312.14870","n_code_links":1,"syntology":null},{"paper":null,"slug":"theory-of-hallucinations-based-on","title":"Theory of Hallucinations based on Equivariance","date":"2023-12-22","arxiv_id":"2312.14504","n_code_links":0,"syntology":null},{"paper":null,"slug":"aspect-based-sentiment-analysis-with-explicit","title":"Aspect-Based Sentiment Analysis with Explicit Sentiment Augmentations","date":"2023-12-18","arxiv_id":"2312.10961","n_code_links":0,"syntology":null},{"paper":"/paper/multilingual-large-language-models-leak-human","slug":"multilingual-large-language-models-leak-human","title":"Multilingual large language models leak human stereotypes across language boundaries","date":"2023-12-12","arxiv_id":"2312.07141","n_code_links":1,"syntology":null},{"paper":null,"slug":"translating-natural-language-queries-to-sql","title":"Translating Natural Language Queries to SQL Using the T5 Model","date":"2023-12-12","arxiv_id":"2312.12414","n_code_links":0,"syntology":null},{"paper":null,"slug":"aikyam-a-video-conferencing-utility-for-deaf","title":"Aikyam: A Video Conferencing Utility for Deaf and Dumb","date":"2023-12-10","arxiv_id":"2312.05962","n_code_links":0,"syntology":null},{"paper":null,"slug":"domain-adaptation-of-a-state-of-the-art-text","title":"Domain Adaptation of a State of the Art Text-to-SQL Model: Lessons Learned and Challenges Found","date":"2023-12-09","arxiv_id":"2312.05448","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-reinforcement-learning-and-large","title":"PerfRL: A Small Language Model Framework for Efficient Code Optimization","date":"2023-12-09","arxiv_id":"2312.05657","n_code_links":0,"syntology":null},{"paper":null,"slug":"converting-epics-stories-into-pseudocode","title":"Converting Epics/Stories into Pseudocode using Transformers","date":"2023-12-08","arxiv_id":"2312.05047","n_code_links":0,"syntology":null},{"paper":null,"slug":"making-translators-privacy-aware-on-the-user","title":"Making Translators Privacy-aware on the User's Side","date":"2023-12-07","arxiv_id":"2312.04068","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-text-to-text-model-for-multilingual","title":"A Text-to-Text Model for Multilingual Offensive Language Identification","date":"2023-12-06","arxiv_id":"2312.03379","n_code_links":0,"syntology":null},{"paper":null,"slug":"compositional-generalization-for-data-to-text","title":"Compositional Generalization for Data-to-Text Generation","date":"2023-12-05","arxiv_id":"2312.02748","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-machine-learning-approach-towards-skill","title":"A Machine Learning Approach Towards SKILL Code Autocompletion","date":"2023-12-04","arxiv_id":"2312.01921","n_code_links":0,"syntology":null},{"paper":"/paper/harnessing-the-power-of-prompt-based","slug":"harnessing-the-power-of-prompt-based","title":"Harnessing the Power of Prompt-based Techniques for Generating School-Level Questions using Large Language Models","date":"2023-12-02","arxiv_id":"2312.01032","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-leveraging-llms-for-conditional-qa","title":"Towards leveraging LLMs for Conditional QA","date":"2023-12-02","arxiv_id":"2312.01143","n_code_links":0,"syntology":null},{"paper":"/paper/data-generation-for-post-ocr-correction-of","slug":"data-generation-for-post-ocr-correction-of","title":"Data Generation for Post-OCR correction of Cyrillic handwriting","date":"2023-11-27","arxiv_id":"2311.15896","n_code_links":2,"syntology":null},{"paper":"/paper/image-super-resolution-with-text-prompt","slug":"image-super-resolution-with-text-prompt","title":"Image Super-Resolution with Text Prompt Diffusion","date":"2023-11-24","arxiv_id":"2311.14282","n_code_links":1,"syntology":null},{"paper":"/paper/compeft-compression-for-communicating","slug":"compeft-compression-for-communicating","title":"ComPEFT: Compression for Communicating Parameter Efficient Updates via Sparsification and Quantization","date":"2023-11-22","arxiv_id":"2311.13171","n_code_links":1,"syntology":{"ran":8,"of":10,"n_ran_checked":7,"n_instrument":1,"unverified":2,"pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["prateeky2806/compeft"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-novel-transformer-based-approach-for-soil","title":"A novel transformer-based approach for soil temperature prediction","date":"2023-11-20","arxiv_id":"2311.11626","n_code_links":0,"syntology":null},{"paper":null,"slug":"bit-cipher-a-simple-yet-powerful-word","title":"Bit Cipher -- A Simple yet Powerful Word Representation System that Integrates Efficiently with Language Models","date":"2023-11-18","arxiv_id":"2311.11012","n_code_links":0,"syntology":null},{"paper":"/paper/vashantor-a-large-scale-multilingual","slug":"vashantor-a-large-scale-multilingual","title":"Vashantor: A Large-scale Multilingual Benchmark Dataset for Automated Translation of Bangla Regional Dialects to Bangla Language","date":"2023-11-18","arxiv_id":"2311.11142","n_code_links":1,"syntology":null},{"paper":"/paper/dynapipe-optimizing-multi-task-training","slug":"dynapipe-optimizing-multi-task-training","title":"DynaPipe: Optimizing Multi-task Training through Dynamic Pipelines","date":"2023-11-17","arxiv_id":"2311.10418","n_code_links":2,"syntology":null},{"paper":null,"slug":"memory-augmented-language-models-through","title":"Memory Augmented Language Models through Mixture of Word Experts","date":"2023-11-15","arxiv_id":"2311.10768","n_code_links":0,"syntology":null},{"paper":"/paper/spot-a-natural-language-interface-for","slug":"spot-a-natural-language-interface-for","title":"Spot: A Natural Language Interface for Geospatial Searches in OSM","date":"2023-11-14","arxiv_id":"2311.08093","n_code_links":1,"syntology":null},{"paper":null,"slug":"ut5-pretraining-non-autoregressive-t5-with","title":"UT5: Pretraining Non autoregressive T5 with unrolled denoising","date":"2023-11-14","arxiv_id":"2311.08552","n_code_links":0,"syntology":null},{"paper":null,"slug":"controllable-topic-focused-abstractive","title":"Controllable Topic-Focused Abstractive Summarization","date":"2023-11-12","arxiv_id":"2311.06724","n_code_links":0,"syntology":null},{"paper":null,"slug":"argumentation-element-annotation-modeling","title":"Argumentation Element Annotation Modeling using XLNet","date":"2023-11-10","arxiv_id":"2311.06239","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-fine-tuning-chatgpt-for-news","title":"Exploring Fine-tuning ChatGPT for News Recommendation","date":"2023-11-10","arxiv_id":"2311.05850","n_code_links":0,"syntology":null},{"paper":"/paper/deep-learning-brasil-at-absapt-2022","slug":"deep-learning-brasil-at-absapt-2022","title":"Deep Learning Brasil at ABSAPT 2022: Portuguese Transformer Ensemble Approaches","date":"2023-11-08","arxiv_id":"2311.05051","n_code_links":1,"syntology":null},{"paper":"/paper/modelling-sentiment-analysis-llms-and-data","slug":"modelling-sentiment-analysis-llms-and-data","title":"Modelling Sentiment Analysis: LLMs and data augmentation techniques","date":"2023-11-07","arxiv_id":"2311.04139","n_code_links":1,"syntology":null},{"paper":null,"slug":"adapting-pre-trained-generative-models-for","title":"Adapting Pre-trained Generative Models for Extractive Question Answering","date":"2023-11-06","arxiv_id":"2311.02961","n_code_links":0,"syntology":null},{"paper":"/paper/famesumm-investigating-and-improving","slug":"famesumm-investigating-and-improving","title":"FaMeSumm: Investigating and Improving Faithfulness of Medical Summarization","date":"2023-11-03","arxiv_id":"2311.02271","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["psunlpgroup/famesumm"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/indo-lego-absa-a-multitask-generative-aspect","slug":"indo-lego-absa-a-multitask-generative-aspect","title":"Indo LEGO-ABSA: A Multitask Generative Aspect Based Sentiment Analysis for Indonesian Language","date":"2023-11-03","arxiv_id":"2311.01757","n_code_links":1,"syntology":null},{"paper":"/paper/better-together-enhancing-generative","slug":"better-together-enhancing-generative","title":"Better Together: Enhancing Generative Knowledge Graph Completion with Language Models and Neighborhood Information","date":"2023-11-02","arxiv_id":"2311.01326","n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-defect-prediction-from-unrealistic","title":"Learning Defect Prediction from Unrealistic Data","date":"2023-11-02","arxiv_id":"2311.00931","n_code_links":0,"syntology":null},{"paper":null,"slug":"maaig-motion-analysis-and-instruction","title":"MAAIG: Motion Analysis And Instruction Generation","date":"2023-11-02","arxiv_id":"2311.00980","n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-alignment-and-flexible-positional","title":"Attention Alignment and Flexible Positional Embeddings Improve Transformer Length Extrapolation","date":"2023-11-01","arxiv_id":"2311.00684","n_code_links":0,"syntology":null},{"paper":"/paper/data-augmentation-for-code-translation-with","slug":"data-augmentation-for-code-translation-with","title":"Data Augmentation for Code Translation with Comparable Corpora and Multiple References","date":"2023-11-01","arxiv_id":"2311.00317","n_code_links":1,"syntology":{"ran":6,"of":12,"n_ran_checked":6,"n_instrument":0,"unverified":6,"pointer_only":12,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 6 unverified","official":{"repos":["veronicium/cmtrans"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":6,"ran_from_kinds":["official"]}}},{"paper":"/paper/generating-medical-instructions-with","slug":"generating-medical-instructions-with","title":"Generating Medical Prescriptions with Conditional Transformer","date":"2023-10-30","arxiv_id":"2310.19727","n_code_links":1,"syntology":{"ran":10,"of":14,"n_ran_checked":7,"n_instrument":3,"unverified":4,"pointer_only":14,"phrase":"10 ran (of which 4 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 3 where Syntology's instrument failed) · 4 unverified","official":{"repos":["hecta-uom/label-to-text-transformer"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":4,"n_ran_no_instrument_failure":7,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"an-ensemble-method-based-on-the-combination","title":"An Ensemble Method Based on the Combination of Transformers with Convolutional Neural Networks to Detect Artificially Generated Text","date":"2023-10-26","arxiv_id":"2310.17312","n_code_links":0,"syntology":null},{"paper":"/paper/lightlm-a-lightweight-deep-and-narrow","slug":"lightlm-a-lightweight-deep-and-narrow","title":"LightLM: A Lightweight Deep and Narrow Language Model for Generative Recommendation","date":"2023-10-26","arxiv_id":"2310.17488","n_code_links":1,"syntology":{"ran":11,"of":12,"n_ran_checked":11,"n_instrument":0,"unverified":1,"pointer_only":12,"phrase":"11 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 1 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["dongyuanjushi/lightlm"],"state":"official (archive's flag): 11 ran","n_ran":11,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/redco-a-lightweight-tool-to-automate","slug":"redco-a-lightweight-tool-to-automate","title":"RedCoast: A Lightweight Tool to Automate Distributed Training of LLMs on Any GPU/TPUs","date":"2023-10-25","arxiv_id":"2310.16355","n_code_links":1,"syntology":{"ran":12,"of":17,"n_ran_checked":12,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["tanyuqian/redco"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"efficient-data-learning-for-open-information","title":"Efficient Data Learning for Open Information Extraction with Pre-trained Language Models","date":"2023-10-23","arxiv_id":"2310.15021","n_code_links":0,"syntology":null},{"paper":"/paper/once-upon-a-textit-time-in-textit-graph","slug":"once-upon-a-textit-time-in-textit-graph","title":"Once Upon a $\\textit{Time}$ in $\\textit{Graph}$: Relative-Time Pretraining for Complex Temporal Reasoning","date":"2023-10-23","arxiv_id":"2310.14709","n_code_links":1,"syntology":{"ran":7,"of":11,"n_ran_checked":7,"n_instrument":0,"unverified":4,"pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["damo-nlp-sg/rememo"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/improving-cross-lingual-transfer-through","slug":"improving-cross-lingual-transfer-through","title":"Improving Cross-Lingual Transfer through Subtree-Aware Word Reordering","date":"2023-10-20","arxiv_id":"2310.13583","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-graph-neural-networks-for-indian","title":"Exploring Graph Neural Networks for Indian Legal Judgment Prediction","date":"2023-10-19","arxiv_id":"2310.12800","n_code_links":0,"syntology":null},{"paper":null,"slug":"empirical-study-of-pretrained-multilingual","title":"Empirical study of pretrained multilingual language models for zero-shot cross-lingual knowledge transfer in generation","date":"2023-10-15","arxiv_id":"2310.09917","n_code_links":0,"syntology":null},{"paper":"/paper/retrieval-generation-alignment-for-end-to-end","slug":"retrieval-generation-alignment-for-end-to-end","title":"Retrieval-Generation Alignment for End-to-End Task-Oriented Dialogue System","date":"2023-10-13","arxiv_id":"2310.08877","n_code_links":1,"syntology":null},{"paper":null,"slug":"answer-candidate-type-selection-text-to-text","title":"Answer Candidate Type Selection: Text-to-Text Language Model for Closed Book Question Answering Meets Knowledge Graphs","date":"2023-10-10","arxiv_id":"2310.07008","n_code_links":0,"syntology":null},{"paper":"/paper/sparse-finetuning-for-inference-acceleration","slug":"sparse-finetuning-for-inference-acceleration","title":"Sparse Fine-tuning for Inference Acceleration of Large Language Models","date":"2023-10-10","arxiv_id":"2310.06927","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["ist-daslab/sparsefinetuning"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"automating-customer-service-using-langchain","title":"Automating Customer Service using LangChain: Building custom open-source GPT Chatbot for organizations","date":"2023-10-09","arxiv_id":"2310.05421","n_code_links":0,"syntology":null},{"paper":null,"slug":"benchmarking-large-language-models-with","title":"Benchmarking Large Language Models with Augmented Instructions for Fine-grained Information Extraction","date":"2023-10-08","arxiv_id":"2310.05092","n_code_links":0,"syntology":null},{"paper":"/paper/slogan-generation-with-noise-perturbation","slug":"slogan-generation-with-noise-perturbation","title":"Effective Slogan Generation with Noise Perturbation","date":"2023-10-06","arxiv_id":"2310.04472","n_code_links":1,"syntology":null},{"paper":"/paper/can-large-language-models-be-good-path","slug":"can-large-language-models-be-good-path","title":"Can Large Language Models be Good Path Planners? A Benchmark and Investigation on Spatial-temporal Reasoning","date":"2023-10-05","arxiv_id":"2310.03249","n_code_links":1,"syntology":{"ran":28,"of":28,"n_ran_checked":28,"n_instrument":0,"unverified":0,"pointer_only":28,"phrase":"28 ran (of which 0 constructed an object rather than computing a result; 28 with no instrument failure: 0 honoured, 0 violated, 28 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["mohamedaghzal/llms-as-path-planners"],"state":"official (archive's flag): 28 ran","n_ran":28,"n_constructed":0,"n_ran_no_instrument_failure":28,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/dspy-compiling-declarative-language-model","slug":"dspy-compiling-declarative-language-model","title":"DSPy: Compiling Declarative Language Model Calls into Self-Improving Pipelines","date":"2023-10-05","arxiv_id":"2310.03714","n_code_links":3,"syntology":{"ran":3,"of":7,"n_ran_checked":3,"n_instrument":0,"unverified":4,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["stanfordnlp/dspy"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"natural-language-models-for-data","title":"Natural Language Models for Data Visualization Utilizing nvBench Dataset","date":"2023-10-02","arxiv_id":"2310.00832","n_code_links":0,"syntology":null},{"paper":null,"slug":"testing-the-limits-of-unified-sequence-to","title":"Testing the Limits of Unified Sequence to Sequence LLM Pretraining on Diverse Table Data Tasks","date":"2023-10-01","arxiv_id":"2310.00789","n_code_links":0,"syntology":null},{"paper":null,"slug":"debertinha-a-multistep-approach-to-adapt","title":"DeBERTinha: A Multistep Approach to Adapt DebertaV3 XSmall for Brazilian Portuguese Natural Language Processing Task","date":"2023-09-28","arxiv_id":"2309.16844","n_code_links":0,"syntology":null},{"paper":"/paper/capp-130-a-corpus-of-chinese-application","slug":"capp-130-a-corpus-of-chinese-application","title":"CAPP-130: A Corpus of Chinese Application Privacy Policy Summarization and Interpretation","date":"2023-09-26","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"program-repair-with-minimal-edits-using","title":"Program Repair with Minimal Edits Using CodeT5","date":"2023-09-26","arxiv_id":"2309.14760","n_code_links":0,"syntology":null},{"paper":"/paper/lexical-squad-multimodal-hate-speech-event","slug":"lexical-squad-multimodal-hate-speech-event","title":"Lexical Squad@Multimodal Hate Speech Event Detection 2023: Multimodal Hate Speech Detection using Fused Ensemble Approach","date":"2023-09-23","arxiv_id":"2309.13354","n_code_links":1,"syntology":null},{"paper":"/paper/on-the-relationship-between-skill-neurons-and","slug":"on-the-relationship-between-skill-neurons-and","title":"On the Relationship between Skill Neurons and Robustness in Prompt Tuning","date":"2023-09-21","arxiv_id":"2309.12263","n_code_links":1,"syntology":null},{"paper":null,"slug":"localize-retrieve-and-fuse-a-generalized","title":"Localize, Retrieve and Fuse: A Generalized Framework for Free-Form Question Answering over Tables","date":"2023-09-20","arxiv_id":"2309.11049","n_code_links":0,"syntology":null},{"paper":"/paper/sequence-to-sequence-spanish-pre-trained","slug":"sequence-to-sequence-spanish-pre-trained","title":"Sequence-to-Sequence Spanish Pre-trained Language Models","date":"2023-09-20","arxiv_id":"2309.11259","n_code_links":1,"syntology":null},{"paper":null,"slug":"accelerating-in-browser-deep-learning","title":"Empowering In-Browser Deep Learning Inference on Edge Devices with Just-in-Time Kernel Optimizations","date":"2023-09-16","arxiv_id":"2309.08978","n_code_links":0,"syntology":null},{"paper":"/paper/structural-self-supervised-objectives-for","slug":"structural-self-supervised-objectives-for","title":"Structural Self-Supervised Objectives for Transformers","date":"2023-09-15","arxiv_id":"2309.08272","n_code_links":1,"syntology":null},{"paper":"/paper/dblplink-an-entity-linker-for-the-dblp","slug":"dblplink-an-entity-linker-for-the-dblp","title":"DBLPLink: An Entity Linker for the DBLP Scholarly Knowledge Graph","date":"2023-09-14","arxiv_id":"2309.07545","n_code_links":1,"syntology":null},{"paper":"/paper/benchmarking-procedural-language","slug":"benchmarking-procedural-language","title":"Benchmarking Procedural Language Understanding for Low-Resource Languages: A Case Study on Turkish","date":"2023-09-13","arxiv_id":"2309.06698","n_code_links":1,"syntology":null},{"paper":null,"slug":"2309-06057","title":"RAP-Gen: Retrieval-Augmented Patch Generation with CodeT5 for Automatic Program Repair","date":"2023-09-12","arxiv_id":"2309.06057","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-natural-language-biases-with-prompt","title":"Detecting Natural Language Biases with Prompt-based Learning","date":"2023-09-11","arxiv_id":"2309.05227","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-nlp-models-identify-distinguish-and","title":"Can NLP Models 'Identify', 'Distinguish', and 'Justify' Questions that Don't have a Definitive Answer?","date":"2023-09-08","arxiv_id":"2309.04635","n_code_links":0,"syntology":null},{"paper":"/paper/nanot5-a-pytorch-framework-for-pre-training","slug":"nanot5-a-pytorch-framework-for-pre-training","title":"nanoT5: A PyTorch Framework for Pre-training and Fine-tuning T5-style Models with Limited Resources","date":"2023-09-05","arxiv_id":"2309.02373","n_code_links":1,"syntology":{"ran":6,"of":9,"n_ran_checked":6,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["piotrnawrot/nanot5"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/transformer-compression-via-subspace","slug":"transformer-compression-via-subspace","title":"$\\rm SP^3$: Enhancing Structured Pruning via PCA Projection","date":"2023-08-31","arxiv_id":"2308.16475","n_code_links":1,"syntology":{"ran":12,"of":14,"n_ran_checked":11,"n_instrument":1,"unverified":2,"pointer_only":14,"phrase":"12 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 1 honoured, 1 violated, 9 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["hyx1999/sp3"],"state":"official (archive's flag): 12 ran","n_ran":12,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/multi-party-goal-tracking-with-llms-comparing","slug":"multi-party-goal-tracking-with-llms-comparing","title":"Multi-party Goal Tracking with LLMs: Comparing Pre-training, Fine-tuning, and Prompt Engineering","date":"2023-08-29","arxiv_id":"2308.15231","n_code_links":1,"syntology":null}],"record_sha256":"3b03f405f5ef88ed81c574a3d1b9252c87417cbce0efc84685f4e01c4a956b85","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}