{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/wordpiece/papers/19","list_of":"/method/wordpiece","method":"WordPiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":19,"pages_in_order":71,"rows_per_page":100,"rows":[1801,1900],"of":7063,"counts":{"archive_papers_tagged":7063,"with_a_code_link":2910,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7063,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":529,"every_run_a_failure_of_syntologys_instrument":121,"listed_with_a_run_with_no_instrument_failure":529,"listed_every_run_a_failure_of_syntologys_instrument":121,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/wordpiece","prev":"/method/wordpiece/papers/18","next":"/method/wordpiece/papers/20","papers":[{"paper":"/paper/exploiting-chatgpt-for-diagnosing-autism","slug":"exploiting-chatgpt-for-diagnosing-autism","title":"Exploiting ChatGPT for Diagnosing Autism-Associated Language Disorders and Identifying Distinct Features","date":"2024-05-03","arxiv_id":"2405.01799","n_code_links":1,"syntology":null},{"paper":null,"slug":"reasons-a-benchmark-for-retrieval-and","title":"Attribution in Scientific Literature: New Benchmark and Methods","date":"2024-05-03","arxiv_id":"2405.02228","n_code_links":0,"syntology":null},{"paper":"/paper/structural-pruning-of-pre-trained-language","slug":"structural-pruning-of-pre-trained-language","title":"Structural Pruning of Pre-trained Language Models via Neural Architecture Search","date":"2024-05-03","arxiv_id":"2405.02267","n_code_links":1,"syntology":null},{"paper":null,"slug":"early-transformers-a-study-on-efficient","title":"Early Transformers: A study on Efficient Training of Transformer Models through Early-Bird Lottery Tickets","date":"2024-05-02","arxiv_id":"2405.02353","n_code_links":0,"syntology":null},{"paper":"/paper/investigating-wit-creativity-and","slug":"investigating-wit-creativity-and","title":"Investigating Wit, Creativity, and Detectability of Large Language Models in Domain-Specific Writing Style Adaptation of Reddit's Showerthoughts","date":"2024-05-02","arxiv_id":"2405.01660","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-named-entity-recognition-and-topic-modeling","title":"A Named Entity Recognition and Topic Modeling-based Solution for Locating and Better Assessment of Natural Disasters in Social Media","date":"2024-05-01","arxiv_id":"2405.00903","n_code_links":0,"syntology":null},{"paper":"/paper/integrating-a-i-in-higher-education-protocol","slug":"integrating-a-i-in-higher-education-protocol","title":"Integrating A.I. in Higher Education: Protocol for a Pilot Study with 'SAMCares: An Adaptive Learning Hub'","date":"2024-05-01","arxiv_id":"2405.00330","n_code_links":1,"syntology":null},{"paper":"/paper/opinion-mining-using-pre-trained-large","slug":"opinion-mining-using-pre-trained-large","title":"Opinion Mining Using Pre-Trained Large Language Models: Identifying the Type, Polarity, Intensity, Expression, and Source of Private States","date":"2024-05-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/graph-neural-network-approach-to-semantic","slug":"graph-neural-network-approach-to-semantic","title":"Graph Neural Network Approach to Semantic Type Detection in Tables","date":"2024-04-30","arxiv_id":"2405.00123","n_code_links":1,"syntology":null},{"paper":"/paper/towards-a-search-engine-for-machines-unified","slug":"towards-a-search-engine-for-machines-unified","title":"Towards a Search Engine for Machines: Unified Ranking for Multiple Retrieval-Augmented Large Language Models","date":"2024-04-30","arxiv_id":"2405.00175","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-cost-minimization-approach-to-fix-the","title":"A cost minimization approach to fix the vocabulary size in a tokenizer for an End-to-End ASR system","date":"2024-04-29","arxiv_id":"2406.02563","n_code_links":0,"syntology":null},{"paper":null,"slug":"federa-efficient-fine-tuning-of-language","title":"FeDeRA:Efficient Fine-tuning of Language Models in Federated Learning Leveraging Weight Decomposition","date":"2024-04-29","arxiv_id":"2404.18848","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-perplexity-predict-fine-tuning","title":"Can Perplexity Predict Fine-Tuning Performance? An Investigation of Tokenization Effects on Sequential Language Models for Nepali","date":"2024-04-28","arxiv_id":"2404.18071","n_code_links":0,"syntology":null},{"paper":"/paper/l3cube-mahanews-news-based-short-text-and","slug":"l3cube-mahanews-news-based-short-text-and","title":"L3Cube-MahaNews: News-based Short Text and Long Document Classification Datasets in Marathi","date":"2024-04-28","arxiv_id":"2404.18216","n_code_links":1,"syntology":null},{"paper":null,"slug":"tabular-embedding-model-tem-finetuning","title":"Tabular Embedding Model (TEM): Finetuning Embedding Models For Tabular RAG Applications","date":"2024-04-28","arxiv_id":"2405.01585","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-pre-trained-generative-language","title":"Enhancing Pre-Trained Generative Language Models with Question Attended Span Extraction on Machine Reading Comprehension","date":"2024-04-27","arxiv_id":"2404.17991","n_code_links":0,"syntology":null},{"paper":null,"slug":"tool-calling-enhancing-medication","title":"Tool Calling: Enhancing Medication Consultation via Retrieval-Augmented Large Language Models","date":"2024-04-27","arxiv_id":"2404.17897","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-legal-compliance-and-regulation","title":"Enhancing Legal Compliance and Regulation Analysis with Large Language Models","date":"2024-04-26","arxiv_id":"2404.17522","n_code_links":0,"syntology":null},{"paper":null,"slug":"human-imperceptible-retrieval-poisoning","title":"Human-Imperceptible Retrieval Poisoning Attacks in LLM-Powered Applications","date":"2024-04-26","arxiv_id":"2404.17196","n_code_links":0,"syntology":null},{"paper":null,"slug":"retrieval-augmented-generation-with-knowledge","title":"Retrieval-Augmented Generation with Knowledge Graphs for Customer Service Question Answering","date":"2024-04-26","arxiv_id":"2404.17723","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-short-survey-of-human-mobility-prediction","title":"A Short Survey of Human Mobility Prediction in Epidemic Modeling from Transformers to LLMs","date":"2024-04-25","arxiv_id":"2404.16921","n_code_links":0,"syntology":null},{"paper":null,"slug":"analise-de-ambiguidade-linguistica-em-modelos","title":"Análise de ambiguidade linguística em modelos de linguagem de grande escala (LLMs)","date":"2024-04-25","arxiv_id":"2404.16653","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-consistency-and-reasoning","title":"Evaluating Consistency and Reasoning Capabilities of Large Language Models","date":"2024-04-25","arxiv_id":"2404.16478","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-internal-numeracy-in-language","title":"Exploring Internal Numeracy in Language Models: A Case Study on ALBERT","date":"2024-04-25","arxiv_id":"2404.16574","n_code_links":0,"syntology":null},{"paper":"/paper/incorporating-lexical-and-syntactic-knowledge","slug":"incorporating-lexical-and-syntactic-knowledge","title":"Incorporating Lexical and Syntactic Knowledge for Unsupervised Cross-Lingual Transfer","date":"2024-04-25","arxiv_id":"2404.16627","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comprehensive-survey-on-evaluating-large","title":"A Comprehensive Survey on Evaluating Large Language Model Applications in the Medical Industry","date":"2024-04-24","arxiv_id":"2404.15777","n_code_links":0,"syntology":null},{"paper":null,"slug":"bert-vs-gpt-for-financial-engineering","title":"BERT vs GPT for financial engineering","date":"2024-04-24","arxiv_id":"2405.12990","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-conceptual-abstraction-in-llms","title":"Detecting Conceptual Abstraction in LLMs","date":"2024-04-24","arxiv_id":"2404.15848","n_code_links":0,"syntology":null},{"paper":"/paper/from-local-to-global-a-graph-rag-approach-to","slug":"from-local-to-global-a-graph-rag-approach-to","title":"From Local to Global: A Graph RAG Approach to Query-Focused Summarization","date":"2024-04-24","arxiv_id":"2404.16130","n_code_links":3,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":null}},{"paper":null,"slug":"investigating-the-prompt-leakage-effect-and","title":"Prompt Leakage effect and defense strategies for multi-turn LLM interactions","date":"2024-04-24","arxiv_id":"2404.16251","n_code_links":0,"syntology":null},{"paper":"/paper/studying-large-language-model-behaviors-under","slug":"studying-large-language-model-behaviors-under","title":"Studying Large Language Model Behaviors Under Context-Memory Conflicts With Real Documents","date":"2024-04-24","arxiv_id":"2404.16032","n_code_links":1,"syntology":{"ran":7,"of":8,"n_ran_checked":7,"n_instrument":0,"unverified":1,"pointer_only":8,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["kortukov/realistic_knowledge_conflicts"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/telco-rag-navigating-the-challenges-of","slug":"telco-rag-navigating-the-challenges-of","title":"Telco-RAG: Navigating the Challenges of Retrieval-Augmented Language Models for Telecommunications","date":"2024-04-24","arxiv_id":"2404.15939","n_code_links":1,"syntology":null},{"paper":null,"slug":"ai-and-machine-learning-for-next-generation","title":"AI and Machine Learning for Next Generation Science Assessments","date":"2024-04-23","arxiv_id":"2405.06660","n_code_links":0,"syntology":null},{"paper":null,"slug":"iryonlp-at-mediqa-corr-2024-tackling-the","title":"IryoNLP at MEDIQA-CORR 2024: Tackling the Medical Error Detection & Correction Task On the Shoulders of Medical Agents","date":"2024-04-23","arxiv_id":"2404.15488","n_code_links":0,"syntology":null},{"paper":"/paper/calc-cmu-at-semeval-2024-task-7-pre-calc","slug":"calc-cmu-at-semeval-2024-task-7-pre-calc","title":"Pre-Calc: Learning to Use the Calculator Improves Numeracy in Language Models","date":"2024-04-22","arxiv_id":"2404.14355","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":0,"n_instrument":3,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","official":{"repos":["calc-cmu/pre-calc"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/llms-know-what-they-need-leveraging-a-missing","slug":"llms-know-what-they-need-leveraging-a-missing","title":"LLMs Know What They Need: Leveraging a Missing Information Guided Framework to Empower Retrieval-Augmented Generation","date":"2024-04-22","arxiv_id":"2404.14043","n_code_links":1,"syntology":null},{"paper":null,"slug":"marking-visual-grading-with-highlighting","title":"Marking: Visual Grading with Highlighting Errors and Annotating Missing Bits","date":"2024-04-22","arxiv_id":"2404.14301","n_code_links":0,"syntology":null},{"paper":"/paper/typos-that-broke-the-rag-s-back-genetic","slug":"typos-that-broke-the-rag-s-back-genetic","title":"Typos that Broke the RAG's Back: Genetic Attack on RAG Pipeline by Simulating Documents in the Wild via Low-level Perturbations","date":"2024-04-22","arxiv_id":"2404.13948","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":5,"n_instrument":0,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["zomss/garag"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/what-do-transformers-know-about-government","slug":"what-do-transformers-know-about-government","title":"What do Transformers Know about Government?","date":"2024-04-22","arxiv_id":"2404.14270","n_code_links":1,"syntology":null},{"paper":"/paper/zero-shot-cross-lingual-stance-detection-via","slug":"zero-shot-cross-lingual-stance-detection-via","title":"Zero-shot Cross-lingual Stance Detection via Adversarial Language Adaptation","date":"2024-04-22","arxiv_id":"2404.14339","n_code_links":1,"syntology":null},{"paper":null,"slug":"automated-text-mining-of-experimental","title":"Automated Text Mining of Experimental Methodologies from Biomedical Literature","date":"2024-04-21","arxiv_id":"2404.13779","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-retrieval-quality-in-retrieval","slug":"evaluating-retrieval-quality-in-retrieval","title":"Evaluating Retrieval Quality in Retrieval-Augmented Generation","date":"2024-04-21","arxiv_id":"2404.13781","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":2,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["alirezasalemi7/erag"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"bert-accelerating-vital-signs-measurement-for","title":"BERT: Accelerating Vital Signs Measurement for Bioradar with An Efficient Recursive Technique","date":"2024-04-20","arxiv_id":"2404.13315","n_code_links":0,"syntology":null},{"paper":"/paper/do-english-named-entity-recognizers-work-well","slug":"do-english-named-entity-recognizers-work-well","title":"Do \"English\" Named Entity Recognizers Work Well on Global Englishes?","date":"2024-04-20","arxiv_id":"2404.13465","n_code_links":2,"syntology":null},{"paper":"/paper/evaluating-subword-tokenization-alien-subword","slug":"evaluating-subword-tokenization-alien-subword","title":"Evaluating Subword Tokenization: Alien Subword Composition and OOV Generalization Challenge","date":"2024-04-20","arxiv_id":"2404.13292","n_code_links":1,"syntology":null},{"paper":"/paper/dubo-sql-diverse-retrieval-augmented","slug":"dubo-sql-diverse-retrieval-augmented","title":"Dubo-SQL: Diverse Retrieval-Augmented Generation and Fine Tuning for Text-to-SQL","date":"2024-04-19","arxiv_id":"2404.12560","n_code_links":1,"syntology":null},{"paper":null,"slug":"enabling-natural-zero-shot-prompting-on","title":"Enabling Natural Zero-Shot Prompting on Encoder Models via Statement-Tuning","date":"2024-04-19","arxiv_id":"2404.12897","n_code_links":0,"syntology":null},{"paper":"/paper/multi-class-depression-detection-through","slug":"multi-class-depression-detection-through","title":"Multi Class Depression Detection Through Tweets using Artificial Intelligence","date":"2024-04-19","arxiv_id":"2404.13104","n_code_links":1,"syntology":null},{"paper":null,"slug":"unlocking-multi-view-insights-in-knowledge","title":"Unlocking Multi-View Insights in Knowledge-Dense Retrieval-Augmented Generation","date":"2024-04-19","arxiv_id":"2404.12879","n_code_links":0,"syntology":null},{"paper":null,"slug":"augmenting-emotion-features-in-irony","title":"Augmenting emotion features in irony detection with Large language modeling","date":"2024-04-18","arxiv_id":"2404.12291","n_code_links":0,"syntology":null},{"paper":null,"slug":"emrqa-msquad-a-medical-dataset-structured","title":"emrQA-msquad: A Medical Dataset Structured with the SQuAD V2.0 Framework, Enriched with emrQA Medical Information","date":"2024-04-18","arxiv_id":"2404.12050","n_code_links":0,"syntology":null},{"paper":null,"slug":"irag-an-incremental-retrieval-augmented","title":"iRAG: Advancing RAG for Videos with an Incremental Approach","date":"2024-04-18","arxiv_id":"2404.12309","n_code_links":0,"syntology":null},{"paper":"/paper/longembed-extending-embedding-models-for-long","slug":"longembed-extending-embedding-models-for-long","title":"LongEmbed: Extending Embedding Models for Long Context Retrieval","date":"2024-04-18","arxiv_id":"2404.12096","n_code_links":1,"syntology":{"ran":5,"of":6,"n_ran_checked":2,"n_instrument":3,"unverified":1,"pointer_only":6,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 1 unverified","official":{"repos":["dwzhu-pku/longembed"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"ragar-your-falsehood-radar-rag-augmented","title":"RAGAR, Your Falsehood Radar: RAG-Augmented Reasoning for Political Fact-Checking using Multimodal Large Language Models","date":"2024-04-18","arxiv_id":"2404.12065","n_code_links":0,"syntology":null},{"paper":null,"slug":"ragcache-efficient-knowledge-caching-for","title":"RAGCache: Efficient Knowledge Caching for Retrieval-Augmented Generation","date":"2024-04-18","arxiv_id":"2404.12457","n_code_links":0,"syntology":null},{"paper":null,"slug":"ram-towards-an-ever-improving-memory-system","title":"RAM: Towards an Ever-Improving Memory System by Learning from Communications","date":"2024-04-18","arxiv_id":"2404.12045","n_code_links":0,"syntology":null},{"paper":null,"slug":"stance-detection-on-social-media-with-fine","title":"Stance Detection on Social Media with Fine-Tuned Large Language Models","date":"2024-04-18","arxiv_id":"2404.12171","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-retrieval-augmented-text-1","title":"A Survey on Retrieval-Augmented Text Generation for Large Language Models","date":"2024-04-17","arxiv_id":"2404.10981","n_code_links":0,"syntology":null},{"paper":null,"slug":"demystifying-legalese-an-automated-approach","title":"Demystifying Legalese: An Automated Approach for Summarizing and Analyzing Overlaps in Privacy Policies and Terms of Service","date":"2024-04-17","arxiv_id":"2404.13087","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-q-a-with-domain-specific-fine","title":"Enhancing Q&A with Domain-Specific Fine-Tuning and Iterative Reasoning: A Comparative Study","date":"2024-04-17","arxiv_id":"2404.11792","n_code_links":0,"syntology":null},{"paper":null,"slug":"improvement-in-semantic-address-matching","title":"Improvement in Semantic Address Matching using Natural Language Processing","date":"2024-04-17","arxiv_id":"2404.11691","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-sentiment-analysis-of-medical-text-based-on","title":"A Sentiment Analysis of Medical Text Based on Deep Learning","date":"2024-04-16","arxiv_id":"2404.10503","n_code_links":0,"syntology":null},{"paper":null,"slug":"bayesjudge-bayesian-kernel-language-modelling","title":"BayesJudge: Bayesian Kernel Language Modelling with Confidence Uncertainty in Legal Judgment Prediction","date":"2024-04-16","arxiv_id":"2404.10481","n_code_links":0,"syntology":null},{"paper":null,"slug":"empowering-interdisciplinary-research-with","title":"Empowering Interdisciplinary Research with BERT-Based Models: An Approach Through SciBERT-CNN with Topic Modeling","date":"2024-04-16","arxiv_id":"2404.13078","n_code_links":0,"syntology":null},{"paper":"/paper/ladic-are-diffusion-models-really-inferior-to","slug":"ladic-are-diffusion-models-really-inferior-to","title":"LaDiC: Are Diffusion Models Really Inferior to Autoregressive Counterparts for Image-to-Text Generation?","date":"2024-04-16","arxiv_id":"2404.10763","n_code_links":1,"syntology":null},{"paper":null,"slug":"relational-graph-convolutional-networks-for-1","title":"Relational Graph Convolutional Networks for Sentiment Analysis","date":"2024-04-16","arxiv_id":"2404.13079","n_code_links":0,"syntology":null},{"paper":"/paper/spiral-of-silences-how-is-large-language","slug":"spiral-of-silences-how-is-large-language","title":"Spiral of Silence: How is Large Language Model Killing Information Retrieval? -- A Case Study on Open Domain Question Answering","date":"2024-04-16","arxiv_id":"2404.10496","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-ai-generated-text-based-on-nlp-and","title":"Detecting AI Generated Text Based on NLP and Machine Learning Approaches","date":"2024-04-15","arxiv_id":"2404.10032","n_code_links":0,"syntology":null},{"paper":"/paper/bert-lsh-reducing-absolute-compute-for","slug":"bert-lsh-reducing-absolute-compute-for","title":"BERT-LSH: Reducing Absolute Compute For Attention","date":"2024-04-12","arxiv_id":"2404.08836","n_code_links":1,"syntology":null},{"paper":null,"slug":"reducing-hallucination-in-structured-outputs","title":"Reducing hallucination in structured outputs via Retrieval-Augmented Generation","date":"2024-04-12","arxiv_id":"2404.08189","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-information-retrieval-evaluation","title":"Generative Information Retrieval Evaluation","date":"2024-04-11","arxiv_id":"2404.08137","n_code_links":0,"syntology":null},{"paper":null,"slug":"emotion-cause-pair-extraction-method-based-on","title":"Emotion-cause pair extraction method based on multi-granularity information and multi-module interaction","date":"2024-04-10","arxiv_id":"2404.06812","n_code_links":0,"syntology":null},{"paper":"/paper/llama-vits-enhancing-tts-synthesis-with","slug":"llama-vits-enhancing-tts-synthesis-with","title":"Llama-VITS: Enhancing TTS Synthesis with Semantic Awareness","date":"2024-04-10","arxiv_id":"2404.06714","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"0 ran · 1 unverified","official":{"repos":["xincanfeng/vitsgpt"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/not-all-contexts-are-equal-teaching-llms","slug":"not-all-contexts-are-equal-teaching-llms","title":"Not All Contexts Are Equal: Teaching LLMs Credibility-aware Generation","date":"2024-04-10","arxiv_id":"2404.06809","n_code_links":1,"syntology":null},{"paper":"/paper/simpler-becomes-harder-do-llms-exhibit-a","slug":"simpler-becomes-harder-do-llms-exhibit-a","title":"Simpler becomes Harder: Do LLMs Exhibit a Coherent Behavior on Simplified Corpora?","date":"2024-04-10","arxiv_id":"2404.06838","n_code_links":1,"syntology":null},{"paper":"/paper/superposition-prompting-improving-and","slug":"superposition-prompting-improving-and","title":"Superposition Prompting: Improving and Accelerating Retrieval-Augmented Generation","date":"2024-04-10","arxiv_id":"2404.06910","n_code_links":1,"syntology":null},{"paper":null,"slug":"comprehensive-study-on-german-language-models","title":"Comprehensive Study on German Language Models for Clinical and Biomedical Text Understanding","date":"2024-04-08","arxiv_id":"2404.05694","n_code_links":0,"syntology":null},{"paper":null,"slug":"medexpqa-multilingual-benchmarking-of-large","title":"MedExpQA: Multilingual Benchmarking of Large Language Models for Medical Question Answering","date":"2024-04-08","arxiv_id":"2404.05590","n_code_links":0,"syntology":null},{"paper":null,"slug":"semantic-stealth-adversarial-text-attacks-on","title":"Semantic Stealth: Adversarial Text Attacks on NLP Using Several Methods","date":"2024-04-08","arxiv_id":"2404.05159","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-morphology-based-investigation-of","title":"A Morphology-Based Investigation of Positional Encodings","date":"2024-04-06","arxiv_id":"2404.04530","n_code_links":0,"syntology":null},{"paper":"/paper/deciphering-political-entity-sentiment-in","slug":"deciphering-political-entity-sentiment-in","title":"Deciphering Political Entity Sentiment in News with Large Language Models: Zero-Shot and Few-Shot Strategies","date":"2024-04-05","arxiv_id":"2404.04361","n_code_links":1,"syntology":null},{"paper":"/paper/investigating-the-robustness-of-modelling","slug":"investigating-the-robustness-of-modelling","title":"Investigating the Robustness of Modelling Decisions for Few-Shot Cross-Topic Stance Detection: A Preregistered Study","date":"2024-04-05","arxiv_id":"2404.03987","n_code_links":1,"syntology":null},{"paper":"/paper/banglaautokg-automatic-bangla-knowledge-graph","slug":"banglaautokg-automatic-bangla-knowledge-graph","title":"BanglaAutoKG: Automatic Bangla Knowledge Graph Construction with Semantic Neural Graph Filtering","date":"2024-04-04","arxiv_id":"2404.03528","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["azminewasi/banglaautokg"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/cbr-rag-case-based-reasoning-for-retrieval","slug":"cbr-rag-case-based-reasoning-for-retrieval","title":"CBR-RAG: Case-Based Reasoning for Retrieval Augmented Generation in LLMs for Legal Question Answering","date":"2024-04-04","arxiv_id":"2404.04302","n_code_links":1,"syntology":null},{"paper":"/paper/conflare-conformal-large-language-model","slug":"conflare-conformal-large-language-model","title":"CONFLARE: CONFormal LArge language model REtrieval","date":"2024-04-04","arxiv_id":"2404.04287","n_code_links":1,"syntology":null},{"paper":"/paper/outlier-efficient-hopfield-layers-for-large","slug":"outlier-efficient-hopfield-layers-for-large","title":"Outlier-Efficient Hopfield Layers for Large Transformer-Based Models","date":"2024-04-04","arxiv_id":"2404.03828","n_code_links":1,"syntology":{"ran":10,"of":13,"n_ran_checked":5,"n_instrument":5,"unverified":3,"pointer_only":0,"phrase":"10 ran (of which 5 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 5 where Syntology's instrument failed) · 3 unverified","official":{"repos":["magics-lab/outeffhop"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":5,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"the-death-of-feature-engineering-bert-with","title":"The Death of Feature Engineering? BERT with Linguistic Features on SQuAD 2.0","date":"2024-04-04","arxiv_id":"2404.03184","n_code_links":0,"syntology":null},{"paper":"/paper/bcamirs-at-semeval-2024-task-4-beyond-words-a","slug":"bcamirs-at-semeval-2024-task-4-beyond-words-a","title":"BCAmirs at SemEval-2024 Task 4: Beyond Words: A Multimodal and Multilingual Exploration of Persuasion in Memes","date":"2024-04-03","arxiv_id":"2404.03022","n_code_links":1,"syntology":null},{"paper":"/paper/clapnq-cohesive-long-form-answers-from","slug":"clapnq-cohesive-long-form-answers-from","title":"CLAPNQ: Cohesive Long-form Answers from Passages in Natural Questions for RAG systems","date":"2024-04-02","arxiv_id":"2404.02103","n_code_links":1,"syntology":null},{"paper":"/paper/ginopic-topic-modeling-with-graph-isomorphism","slug":"ginopic-topic-modeling-with-graph-isomorphism","title":"GINopic: Topic Modeling with Graph Isomorphism Network","date":"2024-04-02","arxiv_id":"2404.02115","n_code_links":1,"syntology":{"ran":10,"of":12,"n_ran_checked":9,"n_instrument":1,"unverified":2,"pointer_only":4,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["adhyasuman/ginopic"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["found_in_text"]}}},{"paper":"/paper/prompts-as-programs-a-structure-aware","slug":"prompts-as-programs-a-structure-aware","title":"Symbolic Prompt Program Search: A Structure-Aware Approach to Efficient Compile-Time Prompt Optimization","date":"2024-04-02","arxiv_id":"2404.02319","n_code_links":1,"syntology":null},{"paper":null,"slug":"advancing-ai-with-integrity-ethical","title":"Advancing AI with Integrity: Ethical Challenges and Solutions in Neural Machine Translation","date":"2024-04-01","arxiv_id":"2404.01070","n_code_links":0,"syntology":null},{"paper":"/paper/aragog-advanced-rag-output-grading","slug":"aragog-advanced-rag-output-grading","title":"ARAGOG: Advanced RAG Output Grading","date":"2024-04-01","arxiv_id":"2404.01037","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-enhanced-retrieval-tool-for-homework","title":"BERT-Enhanced Retrieval Tool for Homework Plagiarism Detection System","date":"2024-04-01","arxiv_id":"2404.01582","n_code_links":0,"syntology":null},{"paper":null,"slug":"securing-social-spaces-harnessing-deep","title":"Securing Social Spaces: Harnessing Deep Learning to Eradicate Cyberbullying","date":"2024-04-01","arxiv_id":"2404.03686","n_code_links":0,"syntology":null},{"paper":null,"slug":"observations-on-building-rag-systems-for","title":"Observations on Building RAG Systems for Technical Documents","date":"2024-03-31","arxiv_id":"2404.00657","n_code_links":0,"syntology":null},{"paper":"/paper/rq-rag-learning-to-refine-queries-for","slug":"rq-rag-learning-to-refine-queries-for","title":"RQ-RAG: Learning to Refine Queries for Retrieval Augmented Generation","date":"2024-03-31","arxiv_id":"2404.00610","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-pre-trained-and-transformer","title":"Leveraging Pre-trained and Transformer-derived Embeddings from EHRs to Characterize Heterogeneity Across Alzheimer's Disease and Related Dementias","date":"2024-03-30","arxiv_id":"2404.00464","n_code_links":0,"syntology":null},{"paper":"/paper/explainable-deep-learning-a-visual-analytics","slug":"explainable-deep-learning-a-visual-analytics","title":"Explainable Deep Learning: A Visual Analytics Approach with Transition Matrices","date":"2024-03-29","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"layernorm-a-key-component-in-parameter","title":"LayerNorm: A key component in parameter-efficient fine-tuning","date":"2024-03-29","arxiv_id":"2403.20284","n_code_links":0,"syntology":null}],"record_sha256":"18cee86df5524a345b29f3807dea026298823f4343e2995a2415a0b8b58d2005","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}