{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/wordpiece/papers/10","list_of":"/method/wordpiece","method":"WordPiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":10,"pages_in_order":71,"rows_per_page":100,"rows":[901,1000],"of":7063,"counts":{"archive_papers_tagged":7063,"with_a_code_link":2910,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7063,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":529,"every_run_a_failure_of_syntologys_instrument":121,"listed_with_a_run_with_no_instrument_failure":529,"listed_every_run_a_failure_of_syntologys_instrument":121,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/wordpiece","prev":"/method/wordpiece/papers/9","next":"/method/wordpiece/papers/11","papers":[{"paper":null,"slug":"scaling-bert-models-for-turkish-automatic","title":"Scaling BERT Models for Turkish Automatic Punctuation and Capitalization Correction","date":"2024-12-03","arxiv_id":"2412.02698","n_code_links":0,"syntology":null},{"paper":null,"slug":"semantic-tokens-in-retrieval-augmented","title":"Semantic Tokens in Retrieval Augmented Generation","date":"2024-12-03","arxiv_id":"2412.02563","n_code_links":0,"syntology":null},{"paper":"/paper/mba-rag-a-bandit-approach-for-adaptive","slug":"mba-rag-a-bandit-approach-for-adaptive","title":"MBA-RAG: a Bandit Approach for Adaptive Retrieval-Augmented Generation through Question Complexity","date":"2024-12-02","arxiv_id":"2412.01572","n_code_links":1,"syntology":{"ran":3,"of":3,"n_ran_checked":1,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["futureeeeee/mba"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"su-roberta-a-semi-supervised-approach-to","title":"Su-RoBERTa: A Semi-supervised Approach to Predicting Suicide Risk through Social Media using Base Language Models","date":"2024-12-02","arxiv_id":"2412.01353","n_code_links":0,"syntology":null},{"paper":"/paper/a-comprehensive-guide-to-explainable-ai-from","slug":"a-comprehensive-guide-to-explainable-ai-from","title":"A Comprehensive Guide to Explainable AI: From Classical Models to LLMs","date":"2024-12-01","arxiv_id":"2412.00800","n_code_links":1,"syntology":null},{"paper":"/paper/lightweight-contenders-navigating-semi","slug":"lightweight-contenders-navigating-semi","title":"Lightweight Contenders: Navigating Semi-Supervised Text Mining through Peer Collaboration and Self Transcendence","date":"2024-12-01","arxiv_id":"2412.00883","n_code_links":1,"syntology":null},{"paper":null,"slug":"does-self-attention-need-separate-weights-in","title":"Does Self-Attention Need Separate Weights in Transformers?","date":"2024-11-30","arxiv_id":"2412.00359","n_code_links":0,"syntology":null},{"paper":null,"slug":"fairness-at-every-intersection-uncovering-and","title":"Fairness at Every Intersection: Uncovering and Mitigating Intersectional Biases in Multimodal Clinical Predictions","date":"2024-11-30","arxiv_id":"2412.00606","n_code_links":0,"syntology":null},{"paper":null,"slug":"advanced-system-integration-analyzing-openapi","title":"Advanced System Integration: Analyzing OpenAPI Chunking for Retrieval-Augmented Generation","date":"2024-11-29","arxiv_id":"2411.19804","n_code_links":0,"syntology":null},{"paper":null,"slug":"generating-a-low-code-complete-workflow-via","title":"Generating a Low-code Complete Workflow via Task Decomposition and RAG","date":"2024-11-29","arxiv_id":"2412.00239","n_code_links":0,"syntology":null},{"paper":null,"slug":"know-your-rag-dataset-taxonomy-and-generation","title":"Know Your RAG: Dataset Taxonomy and Generation Strategies for Evaluating RAG Systems","date":"2024-11-29","arxiv_id":"2411.19710","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-management-for-automobile-failure","title":"Knowledge Management for Automobile Failure Analysis Using Graph RAG","date":"2024-11-29","arxiv_id":"2411.19539","n_code_links":0,"syntology":null},{"paper":null,"slug":"ragdiffusion-faithful-cloth-generation-via","title":"RAGDiffusion: Faithful Cloth Generation via External Knowledge Assimilation","date":"2024-11-29","arxiv_id":"2411.19528","n_code_links":0,"syntology":null},{"paper":null,"slug":"sims-simulating-human-scene-interactions-with","title":"SIMS: Simulating Stylized Human-Scene Interactions with Retrieval-Augmented Script Generation","date":"2024-11-29","arxiv_id":"2411.19921","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-understanding-retrieval-accuracy-and","title":"Towards Understanding Retrieval Accuracy and Prompt Quality in RAG Systems","date":"2024-11-29","arxiv_id":"2411.19463","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficient-learning-content-retrieval-with","title":"Efficient Learning Content Retrieval with Knowledge Injection","date":"2024-11-28","arxiv_id":"2412.00125","n_code_links":0,"syntology":null},{"paper":null,"slug":"knowledge-database-or-poison-base-detecting","title":"RevPRAG: Revealing Poisoning Attacks in Retrieval-Augmented Generation through LLM Activation Analysis","date":"2024-11-28","arxiv_id":"2411.18948","n_code_links":0,"syntology":null},{"paper":"/paper/maskris-semantic-distortion-aware-data","slug":"maskris-semantic-distortion-aware-data","title":"MaskRIS: Semantic Distortion-aware Data Augmentation for Referring Image Segmentation","date":"2024-11-28","arxiv_id":"2411.19067","n_code_links":1,"syntology":null},{"paper":null,"slug":"can-bidirectional-encoder-become-the-ultimate","title":"Can bidirectional encoder become the ultimate winner for downstream applications of foundation models?","date":"2024-11-27","arxiv_id":"2411.18021","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-and-improving-the-robustness-of-1","slug":"evaluating-and-improving-the-robustness-of-1","title":"Evaluating and Improving the Robustness of Security Attack Detectors Generated by LLMs","date":"2024-11-27","arxiv_id":"2411.18216","n_code_links":1,"syntology":null},{"paper":null,"slug":"fine-tuning-large-language-models-for-5","title":"Fine-Tuning Large Language Models for Scientific Text Classification: A Comparative Study","date":"2024-11-27","arxiv_id":"2412.00098","n_code_links":0,"syntology":null},{"paper":null,"slug":"fine-tuning-small-embeddings-for-elevated","title":"Fine-Tuning Small Embeddings for Elevated Performance","date":"2024-11-27","arxiv_id":"2411.18099","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-importance-of-code-mixed-embeddings-for","title":"On Importance of Code-Mixed Embeddings for Hate Speech Identification","date":"2024-11-27","arxiv_id":"2411.18577","n_code_links":0,"syntology":null},{"paper":"/paper/bert-or-fasttext-a-comparative-analysis-of","slug":"bert-or-fasttext-a-comparative-analysis-of","title":"BERT or FastText? A Comparative Analysis of Contextual as well as Non-Contextual Embeddings","date":"2024-11-26","arxiv_id":"2411.17661","n_code_links":1,"syntology":null},{"paper":null,"slug":"fairness-and-performance-in-harmony-data","title":"Fairness And Performance In Harmony: Data Debiasing Is All You Need","date":"2024-11-26","arxiv_id":"2411.17374","n_code_links":0,"syntology":null},{"paper":"/paper/linguistic-laws-meet-protein-sequences-a","slug":"linguistic-laws-meet-protein-sequences-a","title":"Linguistic Laws Meet Protein Sequences: A Comparative Analysis of Subword Tokenization Methods","date":"2024-11-26","arxiv_id":"2411.17669","n_code_links":1,"syntology":null},{"paper":"/paper/what-differentiates-educational-literature-a","slug":"what-differentiates-educational-literature-a","title":"What Differentiates Educational Literature? A Multimodal Fusion Approach of Transformers and Computational Linguistics","date":"2024-11-26","arxiv_id":"2411.17593","n_code_links":0,"syntology":null},{"paper":"/paper/atomr-atomic-operator-empowered-large","slug":"atomr-atomic-operator-empowered-large","title":"AtomR: Atomic Operator-Empowered Large Language Models for Heterogeneous Knowledge Reasoning","date":"2024-11-25","arxiv_id":"2411.16495","n_code_links":1,"syntology":{"ran":13,"of":13,"n_ran_checked":13,"n_instrument":0,"unverified":0,"pointer_only":13,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["THU-KEG/AtomR"],"state":"official (archive's flag): 13 ran","n_ran":13,"n_constructed":0,"n_ran_no_instrument_failure":13,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"dynamic-self-distillation-via-previous-mini","title":"Dynamic Self-Distillation via Previous Mini-batches for Fine-tuning Small Language Models","date":"2024-11-25","arxiv_id":"2411.16991","n_code_links":0,"syntology":null},{"paper":null,"slug":"human-calibrated-automated-testing-and","title":"Human-Calibrated Automated Testing and Validation of Generative Language Models","date":"2024-11-25","arxiv_id":"2411.16391","n_code_links":0,"syntology":null},{"paper":"/paper/lab-rag-label-boosted-retrieval-augmented","slug":"lab-rag-label-boosted-retrieval-augmented","title":"LaB-RAG: Label Boosted Retrieval Augmented Generation for Radiology Report Generation","date":"2024-11-25","arxiv_id":"2411.16523","n_code_links":1,"syntology":null},{"paper":null,"slug":"predictive-power-of-llms-in-financial-markets","title":"Predictive Power of LLMs in Financial Markets","date":"2024-11-25","arxiv_id":"2411.16569","n_code_links":0,"syntology":null},{"paper":null,"slug":"structformer-document-structure-based-masked","title":"StructFormer: Document Structure-based Masked Attention and its Impact on Language Model Pre-Training","date":"2024-11-25","arxiv_id":"2411.16618","n_code_links":0,"syntology":null},{"paper":"/paper/development-of-pre-trained-transformer-based","slug":"development-of-pre-trained-transformer-based","title":"Development of Pre-Trained Transformer-based Models for the Nepali Language","date":"2024-11-24","arxiv_id":"2411.15734","n_code_links":0,"syntology":null},{"paper":null,"slug":"ramie-retrieval-augmented-multi-task","title":"RAMIE: Retrieval-Augmented Multi-task Information Extraction with Large Language Models on Dietary Supplements","date":"2024-11-24","arxiv_id":"2411.15700","n_code_links":0,"syntology":null},{"paper":"/paper/a-comparative-analysis-of-transformer-and","slug":"a-comparative-analysis-of-transformer-and","title":"A Comparative Analysis of Transformer and LSTM Models for Detecting Suicidal Ideation on Reddit","date":"2024-11-23","arxiv_id":"2411.15404","n_code_links":1,"syntology":null},{"paper":null,"slug":"improving-next-tokens-via-second-last","title":"Improving Next Tokens via Second-Last Predictions with Generate and Refine","date":"2024-11-23","arxiv_id":"2411.15661","n_code_links":0,"syntology":null},{"paper":null,"slug":"inducing-human-like-biases-in-moral-reasoning","title":"Inducing Human-like Biases in Moral Reasoning Language Models","date":"2024-11-23","arxiv_id":"2411.15386","n_code_links":0,"syntology":null},{"paper":null,"slug":"traditional-chinese-medicine-case-analysis","title":"Traditional Chinese Medicine Case Analysis System for High-Level Semantic Abstraction: Optimized with Prompt and RAG","date":"2024-11-23","arxiv_id":"2411.15491","n_code_links":0,"syntology":null},{"paper":null,"slug":"astro-hep-bert-a-bidirectional-language-model","title":"Astro-HEP-BERT: A bidirectional language model for studying the meanings of concepts in astrophysics and high energy physics","date":"2024-11-22","arxiv_id":"2411.14877","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-pooling-mechanisms-in","title":"Comparative Analysis of Pooling Mechanisms in LLMs: A Sentiment Analysis Perspective","date":"2024-11-22","arxiv_id":"2411.14654","n_code_links":0,"syntology":null},{"paper":"/paper/kbada-efficient-self-adaptation-on-specific","slug":"kbada-efficient-self-adaptation-on-specific","title":"KBAlign: Efficient Self Adaptation on Specific Knowledge Bases","date":"2024-11-22","arxiv_id":"2411.14790","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-experimental-study-on-data-augmentation","title":"An Experimental Study on Data Augmentation Techniques for Named Entity Recognition on Low-Resource Domains","date":"2024-11-21","arxiv_id":"2411.14551","n_code_links":0,"syntology":null},{"paper":"/paper/bert-based-approach-for-automating-course","slug":"bert-based-approach-for-automating-course","title":"BERT-Based Approach for Automating Course Articulation Matrix Construction with Explainable AI","date":"2024-11-21","arxiv_id":"2411.14254","n_code_links":1,"syntology":null},{"paper":null,"slug":"fastrag-retrieval-augmented-generation-for","title":"FastRAG: Retrieval Augmented Generation for Semi-structured Data","date":"2024-11-21","arxiv_id":"2411.13773","n_code_links":0,"syntology":null},{"paper":"/paper/g-rag-knowledge-expansion-in-material-science","slug":"g-rag-knowledge-expansion-in-material-science","title":"G-RAG: Knowledge Expansion in Material Science","date":"2024-11-21","arxiv_id":"2411.14592","n_code_links":1,"syntology":null},{"paper":"/paper/pos-tagging-to-highlight-the-skeletal","slug":"pos-tagging-to-highlight-the-skeletal","title":"POS-tagging to highlight the skeletal structure of sentences","date":"2024-11-21","arxiv_id":"2411.14393","n_code_links":2,"syntology":null},{"paper":null,"slug":"towards-knowledge-checking-in-retrieval","title":"Towards Knowledge Checking in Retrieval-augmented Generation: A Representation Perspective","date":"2024-11-21","arxiv_id":"2411.14572","n_code_links":0,"syntology":null},{"paper":"/paper/combining-autoregressive-and-autoencoder","slug":"combining-autoregressive-and-autoencoder","title":"Combining Autoregressive and Autoencoder Language Models for Text Classification","date":"2024-11-20","arxiv_id":"2411.13282","n_code_links":1,"syntology":null},{"paper":null,"slug":"dmqr-rag-diverse-multi-query-rewriting-for","title":"DMQR-RAG: Diverse Multi-Query Rewriting for RAG","date":"2024-11-20","arxiv_id":"2411.13154","n_code_links":0,"syntology":null},{"paper":null,"slug":"multimodal-large-language-model-for-wheat","title":"Multimodal large language model for wheat breeding: a new exploration of smart breeding","date":"2024-11-20","arxiv_id":"2411.15203","n_code_links":0,"syntology":null},{"paper":null,"slug":"on-the-way-to-llm-personalization-learning-to","title":"On the Way to LLM Personalization: Learning to Remember User Conversations","date":"2024-11-20","arxiv_id":"2411.13405","n_code_links":0,"syntology":null},{"paper":null,"slug":"retrieval-augmented-generation-for-domain-1","title":"Retrieval-Augmented Generation for Domain-Specific Question Answering: A Case Study on Pittsburgh and CMU","date":"2024-11-20","arxiv_id":"2411.13691","n_code_links":0,"syntology":null},{"paper":null,"slug":"unlocking-historical-clinical-trial-data-with","title":"Unlocking Historical Clinical Trial Data with ALIGN: A Compositional Large Language Model System for Medical Coding","date":"2024-11-20","arxiv_id":"2411.13163","n_code_links":0,"syntology":null},{"paper":"/paper/dlbacktrace-a-model-agnostic-explainability","slug":"dlbacktrace-a-model-agnostic-explainability","title":"DLBacktrace: A Model Agnostic Explainability for any Deep Learning Models","date":"2024-11-19","arxiv_id":"2411.12643","n_code_links":1,"syntology":null},{"paper":null,"slug":"enhancing-multi-class-disease-classification","title":"Enhancing Multi-Class Disease Classification: Neoplasms, Cardiovascular, Nervous System, and Digestive Disorders Using Advanced LLMs","date":"2024-11-19","arxiv_id":"2411.12712","n_code_links":0,"syntology":null},{"paper":null,"slug":"strengthening-fake-news-detection-leveraging","title":"Strengthening Fake News Detection: Leveraging SVM and Sophisticated Text Vectorization Techniques. Defying BERT?","date":"2024-11-19","arxiv_id":"2411.12703","n_code_links":0,"syntology":null},{"paper":"/paper/cnmbert-a-model-for-hanyu-pinyin-abbreviation","slug":"cnmbert-a-model-for-hanyu-pinyin-abbreviation","title":"CNMBERT: A Model for Converting Hanyu Pinyin Abbreviations to Chinese Characters","date":"2024-11-18","arxiv_id":"2411.11770","n_code_links":1,"syntology":null},{"paper":null,"slug":"suicide-risk-assessment-on-social-media-with","title":"Suicide Risk Assessment on Social Media with Semi-Supervised Learning","date":"2024-11-18","arxiv_id":"2411.12767","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-student-sentiment-on-mental","title":"Understanding Student Sentiment on Mental Health Support in Colleges Using Large Language Models","date":"2024-11-18","arxiv_id":"2412.04326","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-novel-approach-to-eliminating","title":"A Novel Approach to Eliminating Hallucinations in Large Language Model-Assisted Causal Discovery","date":"2024-11-16","arxiv_id":"2411.12759","n_code_links":0,"syntology":null},{"paper":null,"slug":"debias-clr-a-contrastive-learning-based","title":"Debias-CLR: A Contrastive Learning Based Debiasing Method for Algorithmic Fairness in Healthcare Applications","date":"2024-11-15","arxiv_id":"2411.10544","n_code_links":0,"syntology":null},{"paper":"/paper/hysteresis-activation-function-for-efficient","slug":"hysteresis-activation-function-for-efficient","title":"Hysteresis Activation Function for Efficient Inference","date":"2024-11-15","arxiv_id":"2411.10573","n_code_links":1,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"0 ran · 1 unverified","official":{"repos":["idankdev/helu"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/information-extraction-from-clinical-notes","slug":"information-extraction-from-clinical-notes","title":"Information Extraction from Clinical Notes: Are We Ready to Switch to Large Language Models?","date":"2024-11-15","arxiv_id":"2411.10020","n_code_links":1,"syntology":null},{"paper":null,"slug":"softlms-efficient-adaptive-low-rank","title":"SoftLMs: Efficient Adaptive Low-Rank Approximation of Language Models using Soft-Thresholding Mechanism","date":"2024-11-15","arxiv_id":"2411.10543","n_code_links":0,"syntology":null},{"paper":null,"slug":"adopting-rag-for-llm-aided-future-vehicle","title":"Adopting RAG for LLM-Aided Future Vehicle Design","date":"2024-11-14","arxiv_id":"2411.09590","n_code_links":0,"syntology":null},{"paper":null,"slug":"comprehensive-and-practical-evaluation-of","title":"Comprehensive and Practical Evaluation of Retrieval-Augmented Generation Systems for Medical Question Answering","date":"2024-11-14","arxiv_id":"2411.09213","n_code_links":0,"syntology":null},{"paper":"/paper/initial-nugget-evaluation-results-for-the","slug":"initial-nugget-evaluation-results-for-the","title":"Initial Nugget Evaluation Results for the TREC 2024 RAG Track with the AutoNuggetizer Framework","date":"2024-11-14","arxiv_id":"2411.09607","n_code_links":2,"syntology":null},{"paper":"/paper/a-large-scale-study-of-relevance-assessments","slug":"a-large-scale-study-of-relevance-assessments","title":"A Large-Scale Study of Relevance Assessments with Large Language Models: An Initial Look","date":"2024-11-13","arxiv_id":"2411.08275","n_code_links":1,"syntology":null},{"paper":null,"slug":"analyst-reports-and-stock-performance","title":"Analyst Reports and Stock Performance: Evidence from the Chinese Market","date":"2024-11-13","arxiv_id":"2411.08726","n_code_links":0,"syntology":null},{"paper":null,"slug":"camembert-2-0-a-smarter-french-language-model","title":"CamemBERT 2.0: A Smarter French Language Model Aged to Perfection","date":"2024-11-13","arxiv_id":"2411.08868","n_code_links":0,"syntology":null},{"paper":"/paper/logllm-log-based-anomaly-detection-using","slug":"logllm-log-based-anomaly-detection-using","title":"LogLLM: Log-based Anomaly Detection Using Large Language Models","date":"2024-11-13","arxiv_id":"2411.08561","n_code_links":1,"syntology":null},{"paper":null,"slug":"towards-optimizing-a-retrieval-augmented","title":"Towards Optimizing a Retrieval Augmented Generation using Large Language Model on Academic Data","date":"2024-11-13","arxiv_id":"2411.08438","n_code_links":0,"syntology":null},{"paper":"/paper/controlled-evaluation-of-syntactic-knowledge","slug":"controlled-evaluation-of-syntactic-knowledge","title":"Controlled Evaluation of Syntactic Knowledge in Multilingual Language Models","date":"2024-11-12","arxiv_id":"2411.07474","n_code_links":1,"syntology":null},{"paper":null,"slug":"query-optimization-for-parametric-knowledge","title":"Query Optimization for Parametric Knowledge Refinement in Retrieval-Augmented Large Language Models","date":"2024-11-12","arxiv_id":"2411.07820","n_code_links":0,"syntology":null},{"paper":"/paper/retrieval-augmented-time-series-forecasting","slug":"retrieval-augmented-time-series-forecasting","title":"Retrieval Augmented Time Series Forecasting","date":"2024-11-12","arxiv_id":"2411.08249","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["kutaytire/retrieval-augmented-time-series-forecasting"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"trustful-llms-customizing-and-grounding-text","title":"Trustful LLMs: Customizing and Grounding Text Generation with Knowledge Bases and Dual Decoders","date":"2024-11-12","arxiv_id":"2411.07870","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-unified-multi-task-learning-architecture","title":"A Unified Multi-Task Learning Architecture for Hate Detection Leveraging User-Based Information","date":"2024-11-11","arxiv_id":"2411.06855","n_code_links":0,"syntology":null},{"paper":"/paper/assistrag-boosting-the-potential-of-large","slug":"assistrag-boosting-the-potential-of-large","title":"AssistRAG: Boosting the Potential of Large Language Models with an Intelligent Information Assistant","date":"2024-11-11","arxiv_id":"2411.06805","n_code_links":1,"syntology":{"ran":10,"of":13,"n_ran_checked":8,"n_instrument":2,"unverified":3,"pointer_only":13,"phrase":"10 ran (of which 1 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 0 violated, 8 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["smallporridge/assistrag"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":1,"n_ran_no_instrument_failure":8,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/autonomous-droplet-microfluidic-design","slug":"autonomous-droplet-microfluidic-design","title":"Autonomous Droplet Microfluidic Design Framework with Large Language Models","date":"2024-11-11","arxiv_id":"2411.06691","n_code_links":1,"syntology":null},{"paper":null,"slug":"evaluating-large-language-models-on-financial","title":"Evaluating Large Language Models on Financial Report Summarization: An Empirical Study","date":"2024-11-11","arxiv_id":"2411.06852","n_code_links":0,"syntology":null},{"paper":null,"slug":"invar-rag-invariant-llm-aligned-retrieval-for","title":"Invar-RAG: Invariant LLM-aligned Retrieval for Better Generation","date":"2024-11-11","arxiv_id":"2411.07021","n_code_links":0,"syntology":null},{"paper":null,"slug":"la4sr-illuminating-the-dark-proteome-with","title":"LA4SR: illuminating the dark proteome with generative AI","date":"2024-11-11","arxiv_id":"2411.06798","n_code_links":0,"syntology":null},{"paper":null,"slug":"tempcharbert-keystroke-dynamics-for","title":"TempCharBERT: Keystroke Dynamics for Continuous Access Control Based on Pre-trained Language Models","date":"2024-11-11","arxiv_id":"2411.07224","n_code_links":0,"syntology":null},{"paper":null,"slug":"token2wave","title":"The Backpropagation of the Wave Network","date":"2024-11-11","arxiv_id":"2411.06989","n_code_links":0,"syntology":null},{"paper":"/paper/toward-optimal-search-and-retrieval-for-rag","slug":"toward-optimal-search-and-retrieval-for-rag","title":"Toward Optimal Search and Retrieval for RAG","date":"2024-11-11","arxiv_id":"2411.07396","n_code_links":1,"syntology":null},{"paper":"/paper/region-aware-text-to-image-generation-via","slug":"region-aware-text-to-image-generation-via","title":"Region-Aware Text-to-Image Generation via Hard Binding and Soft Refinement","date":"2024-11-10","arxiv_id":"2411.06558","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":1,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["nju-pcalab/rag-diffusion"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official","unlocated"]}}},{"paper":null,"slug":"clustering-algorithms-and-rag-enhancing-semi","title":"Clustering Algorithms and RAG Enhancing Semi-Supervised Text Classification with Large LLMs","date":"2024-11-09","arxiv_id":"2411.06175","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-knowledge-boundaries-in-large","title":"Exploring Knowledge Boundaries in Large Language Models for Retrieval Judgment","date":"2024-11-09","arxiv_id":"2411.06207","n_code_links":0,"syntology":null},{"paper":null,"slug":"improved-intent-classification-based-on","title":"Improved intent classification based on context information using a windows-based approach","date":"2024-11-09","arxiv_id":"2411.06022","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-retrieval-augmented-generation-for-1","title":"Leveraging Retrieval-Augmented Generation for Persian University Knowledge Retrieval","date":"2024-11-09","arxiv_id":"2411.06237","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-detection-of-llm-generated-text-a","title":"Robust Detection of LLM-Generated Text: A Comparative Analysis","date":"2024-11-09","arxiv_id":"2411.06248","n_code_links":0,"syntology":null},{"paper":null,"slug":"sufficient-context-a-new-lens-on-retrieval","title":"Sufficient Context: A New Lens on Retrieval Augmented Generation Systems","date":"2024-11-09","arxiv_id":"2411.06037","n_code_links":0,"syntology":null},{"paper":"/paper/a-taxonomy-of-agentops-for-enabling","slug":"a-taxonomy-of-agentops-for-enabling","title":"AgentOps: Enabling Observability of LLM Agents","date":"2024-11-08","arxiv_id":"2411.05285","n_code_links":1,"syntology":null},{"paper":"/paper/findver-explainable-claim-verification-over","slug":"findver-explainable-claim-verification-over","title":"FinDVer: Explainable Claim Verification over Long and Hybrid-Content Financial Documents","date":"2024-11-08","arxiv_id":"2411.05764","n_code_links":1,"syntology":null},{"paper":"/paper/heartbert-a-self-supervised-ecg-embedding","slug":"heartbert-a-self-supervised-ecg-embedding","title":"HeartBERT: A Self-Supervised ECG Embedding Model for Efficient and Effective Medical Signal Analysis","date":"2024-11-08","arxiv_id":"2411.11896","n_code_links":1,"syntology":null},{"paper":"/paper/intellbot-retrieval-augmented-llm-chatbot-for","slug":"intellbot-retrieval-augmented-llm-chatbot-for","title":"IntellBot: Retrieval Augmented LLM Chatbot for Cyber Threat Knowledge Delivery","date":"2024-11-08","arxiv_id":"2411.05442","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-document-financial-question-answering","title":"Multi-Document Financial Question Answering using LLMs","date":"2024-11-08","arxiv_id":"2411.07264","n_code_links":0,"syntology":null},{"paper":null,"slug":"qwen2-5-32b-leveraging-self-consistent-tool","title":"Qwen2.5-32B: Leveraging Self-Consistent Tool-Integrated Reasoning for Bengali Mathematical Olympiad Problem Solving","date":"2024-11-08","arxiv_id":"2411.05934","n_code_links":0,"syntology":null},{"paper":"/paper/sentiment-analysis-of-cyberbullying-data-in","slug":"sentiment-analysis-of-cyberbullying-data-in","title":"Sentiment Analysis of Cyberbullying Data in Social Media","date":"2024-11-08","arxiv_id":"2411.05958","n_code_links":1,"syntology":null}],"record_sha256":"b3053512dd46189fb5c6bb5926b9db15920d6cce343d070c00d87d1f3283a6aa","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}