{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/wordpiece/papers/18","list_of":"/method/wordpiece","method":"WordPiece","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":18,"pages_in_order":71,"rows_per_page":100,"rows":[1701,1800],"of":7063,"counts":{"archive_papers_tagged":7063,"with_a_code_link":2910,"where_syntology_ran_a_sample":650,"not_listed_spam_title":0,"listed":7063,"listed_where_code_ran":650,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":529,"every_run_a_failure_of_syntologys_instrument":121,"listed_with_a_run_with_no_instrument_failure":529,"listed_every_run_a_failure_of_syntologys_instrument":121,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/wordpiece","prev":"/method/wordpiece/papers/17","next":"/method/wordpiece/papers/19","papers":[{"paper":"/paper/mix-of-granularity-optimize-the-chunking","slug":"mix-of-granularity-optimize-the-chunking","title":"Mix-of-Granularity: Optimize the Chunking Granularity for Retrieval-Augmented Generation","date":"2024-06-01","arxiv_id":"2406.00456","n_code_links":1,"syntology":null},{"paper":null,"slug":"pseudo-label-based-domain-adaptation-for-zero","title":"Pseudo-label Based Domain Adaptation for Zero-Shot Text Steganalysis","date":"2024-06-01","arxiv_id":"2406.18565","n_code_links":0,"syntology":null},{"paper":"/paper/roberta-bilstm-a-context-aware-hybrid-model","slug":"roberta-bilstm-a-context-aware-hybrid-model","title":"RoBERTa-BiLSTM: A Context-Aware Hybrid Model for Sentiment Analysis","date":"2024-06-01","arxiv_id":"2406.00367","n_code_links":1,"syntology":null},{"paper":"/paper/a-comparison-of-correspondence-analysis-with","slug":"a-comparison-of-correspondence-analysis-with","title":"A comparison of correspondence analysis with PMI-based word embedding methods","date":"2024-05-31","arxiv_id":"2405.20895","n_code_links":1,"syntology":null},{"paper":null,"slug":"bi-directional-transformers-vs-word2vec","title":"Bi-Directional Transformers vs. word2vec: Discovering Vulnerabilities in Lifted Compiled Code","date":"2024-05-31","arxiv_id":"2405.20611","n_code_links":0,"syntology":null},{"paper":null,"slug":"effect-of-antibody-levels-on-the-spread-of","title":"Effect of antibody levels on the spread of disease in multiple infections","date":"2024-05-31","arxiv_id":"2405.20702","n_code_links":0,"syntology":null},{"paper":"/paper/enhancing-noise-robustness-of-retrieval","slug":"enhancing-noise-robustness-of-retrieval","title":"Enhancing Noise Robustness of Retrieval-Augmented Language Models with Adaptive Adversarial Training","date":"2024-05-31","arxiv_id":"2405.20978","n_code_links":1,"syntology":{"ran":9,"of":9,"n_ran_checked":9,"n_instrument":0,"unverified":0,"pointer_only":9,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["calubkk/raat"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"rag-does-not-work-for-enterprises","title":"RAG Does Not Work for Enterprises","date":"2024-05-31","arxiv_id":"2406.04369","n_code_links":0,"syntology":null},{"paper":null,"slug":"retrieval-meets-reasoning-even-high-school","title":"Retrieval Meets Reasoning: Even High-school Textbook Knowledge Benefits Multimodal Reasoning","date":"2024-05-31","arxiv_id":"2405.20834","n_code_links":0,"syntology":null},{"paper":null,"slug":"ensemble-model-with-bert-roberta-and-xlnet","title":"Ensemble Model With Bert,Roberta and Xlnet For Molecular property prediction","date":"2024-05-30","arxiv_id":"2406.06553","n_code_links":0,"syntology":null},{"paper":"/paper/gnn-rag-graph-neural-retrieval-for-large","slug":"gnn-rag-graph-neural-retrieval-for-large","title":"GNN-RAG: Graph Neural Retrieval for Large Language Model Reasoning","date":"2024-05-30","arxiv_id":"2405.20139","n_code_links":1,"syntology":{"ran":6,"of":7,"n_ran_checked":6,"n_instrument":0,"unverified":1,"pointer_only":7,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["cmavro/gnn-rag"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"is-my-data-in-your-retrieval-database","title":"Is My Data in Your Retrieval Database? Membership Inference Attacks Against Retrieval Augmented Generation","date":"2024-05-30","arxiv_id":"2405.20446","n_code_links":0,"syntology":null},{"paper":null,"slug":"kerascv-and-kerasnlp-vision-and-language","title":"KerasCV and KerasNLP: Vision and Language Power-Ups","date":"2024-05-30","arxiv_id":"2405.20247","n_code_links":0,"syntology":null},{"paper":"/paper/one-token-can-help-learning-scalable-and","slug":"one-token-can-help-learning-scalable-and","title":"One Token Can Help! Learning Scalable and Pluggable Virtual Tokens for Retrieval-Augmented Large Language Models","date":"2024-05-30","arxiv_id":"2405.19670","n_code_links":2,"syntology":null},{"paper":null,"slug":"phantom-general-trigger-attacks-on-retrieval","title":"Phantom: General Trigger Attacks on Retrieval Augmented Language Generation","date":"2024-05-30","arxiv_id":"2405.20485","n_code_links":0,"syntology":null},{"paper":"/paper/student-answer-forecasting-transformer-driven","slug":"student-answer-forecasting-transformer-driven","title":"Student Answer Forecasting: Transformer-Driven Answer Choice Prediction for Language Learning","date":"2024-05-30","arxiv_id":"2405.20079","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-multi-source-retrieval-question-answering","title":"A Multi-Source Retrieval Question Answering Framework Based on RAG","date":"2024-05-29","arxiv_id":"2405.19207","n_code_links":0,"syntology":null},{"paper":null,"slug":"can-gpt-redefine-medical-understanding","title":"Can GPT Redefine Medical Understanding? Evaluating GPT on Biomedical Machine Reading Comprehension","date":"2024-05-29","arxiv_id":"2405.18682","n_code_links":0,"syntology":null},{"paper":"/paper/ctrla-adaptive-retrieval-augmented-generation","slug":"ctrla-adaptive-retrieval-augmented-generation","title":"CtrlA: Adaptive Retrieval-Augmented Generation via Inherent Control","date":"2024-05-29","arxiv_id":"2405.18727","n_code_links":1,"syntology":{"ran":7,"of":12,"n_ran_checked":7,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 0 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["hsliu-initial/ctrla"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"stat-shrinking-transformers-after-training","title":"STAT: Shrinking Transformers After Training","date":"2024-05-29","arxiv_id":"2406.00061","n_code_links":0,"syntology":null},{"paper":"/paper/toward-conversational-agents-with-context-and","slug":"toward-conversational-agents-with-context-and","title":"Toward Conversational Agents with Context and Time Sensitive Long-term Memory","date":"2024-05-29","arxiv_id":"2406.00057","n_code_links":1,"syntology":null},{"paper":null,"slug":"two-layer-retrieval-augmented-generation","title":"Two-Layer Retrieval-Augmented Generation Framework for Low-Resource Medical Question Answering Using Reddit Data: Proof-of-Concept Study","date":"2024-05-29","arxiv_id":"2405.19519","n_code_links":0,"syntology":null},{"paper":"/paper/atm-adversarial-tuning-multi-agent-system","slug":"atm-adversarial-tuning-multi-agent-system","title":"ATM: Adversarial Tuning Multi-agent System Makes a Robust Retrieval-Augmented Generator","date":"2024-05-28","arxiv_id":"2405.18111","n_code_links":1,"syntology":{"ran":6,"of":10,"n_ran_checked":6,"n_instrument":0,"unverified":4,"pointer_only":10,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["chuhac/atm-rag"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"attention-based-sequential-recommendation","title":"Attention-based sequential recommendation system using multimodal data","date":"2024-05-28","arxiv_id":"2405.17959","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-forget-to-connect-improving-rag-with","title":"Don't Forget to Connect! Improving RAG with Graph-based Reranking","date":"2024-05-28","arxiv_id":"2405.18414","n_code_links":0,"syntology":null},{"paper":null,"slug":"widin-wording-image-for-domain-invariant","title":"WIDIn: Wording Image for Domain-Invariant Representation in Single-Source Domain Generalization","date":"2024-05-28","arxiv_id":"2405.18405","n_code_links":0,"syntology":null},{"paper":null,"slug":"augmenting-textual-generation-via-topology","title":"Augmenting Textual Generation via Topology Aware Retrieval","date":"2024-05-27","arxiv_id":"2405.17602","n_code_links":0,"syntology":null},{"paper":"/paper/deeperimpact-optimizing-sparse-learned-index","slug":"deeperimpact-optimizing-sparse-learned-index","title":"DeeperImpact: Optimizing Sparse Learned Index Structures","date":"2024-05-27","arxiv_id":"2405.17093","n_code_links":1,"syntology":null},{"paper":null,"slug":"detecting-deceptive-dark-patterns-in-e","title":"Detecting Deceptive Dark Patterns in E-commerce Platforms","date":"2024-05-27","arxiv_id":"2406.01608","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploiting-the-layered-intrinsic","title":"Exploiting the Layered Intrinsic Dimensionality of Deep Models for Practical Adversarial Training","date":"2024-05-27","arxiv_id":"2405.17130","n_code_links":0,"syntology":null},{"paper":null,"slug":"nv-embed-improved-techniques-for-training","title":"NV-Embed: Improved Techniques for Training LLMs as Generalist Embedding Models","date":"2024-05-27","arxiv_id":"2405.17428","n_code_links":0,"syntology":null},{"paper":null,"slug":"pae-llm-based-product-attribute-extraction","title":"PAE: LLM-based Product Attribute Extraction for E-Commerce Fashion Trends","date":"2024-05-27","arxiv_id":"2405.17533","n_code_links":0,"syntology":null},{"paper":null,"slug":"performance-evaluation-of-reddit-comments","title":"Performance evaluation of Reddit Comments using Machine Learning and Natural Language Processing methods in Sentiment Analysis","date":"2024-05-27","arxiv_id":"2405.16810","n_code_links":0,"syntology":null},{"paper":null,"slug":"qub-cirdan-at-discharge-me-zero-shot","title":"QUB-Cirdan at \"Discharge Me!\": Zero shot discharge letter generation by open-source LLM","date":"2024-05-27","arxiv_id":"2406.00041","n_code_links":0,"syntology":null},{"paper":"/paper/video-enriched-retrieval-augmented-generation","slug":"video-enriched-retrieval-augmented-generation","title":"Video Enriched Retrieval Augmented Generation Using Aligned Video Captions","date":"2024-05-27","arxiv_id":"2405.17706","n_code_links":1,"syntology":null},{"paper":null,"slug":"ai-generated-text-detection-and","title":"AI-Generated Text Detection and Classification Based on BERT Deep Learning Algorithm","date":"2024-05-26","arxiv_id":"2405.16422","n_code_links":0,"syntology":null},{"paper":"/paper/grag-graph-retrieval-augmented-generation","slug":"grag-graph-retrieval-augmented-generation","title":"GRAG: Graph Retrieval-Augmented Generation","date":"2024-05-26","arxiv_id":"2405.16506","n_code_links":1,"syntology":null},{"paper":null,"slug":"m-rag-reinforcing-large-language-model","title":"M-RAG: Reinforcing Large Language Model Performance through Retrieval-Augmented Generation with Multiple Partitions","date":"2024-05-26","arxiv_id":"2405.16420","n_code_links":0,"syntology":null},{"paper":null,"slug":"accelerating-inference-of-retrieval-augmented","title":"Accelerating Inference of Retrieval-Augmented Generation via Sparse Context Selection","date":"2024-05-25","arxiv_id":"2405.16178","n_code_links":0,"syntology":null},{"paper":null,"slug":"incremental-comprehension-of-garden-path","title":"Incremental Comprehension of Garden-Path Sentences by Large Language Models: Semantic Interpretation, Syntactic Re-Analysis, and Attention","date":"2024-05-25","arxiv_id":"2405.16042","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-unlocking-insights-from-logbooks","title":"Towards Unlocking Insights from Logbooks Using AI","date":"2024-05-25","arxiv_id":"2406.12881","n_code_links":0,"syntology":null},{"paper":null,"slug":"enhancing-augmentative-and-alternative","title":"Enhancing Augmentative and Alternative Communication with Card Prediction and Colourful Semantics","date":"2024-05-24","arxiv_id":"2405.15896","n_code_links":0,"syntology":null},{"paper":null,"slug":"textit-comet-a-underline-com-munication","title":"Comet: A Communication-efficient and Performant Approximation for Private Transformer Inference","date":"2024-05-24","arxiv_id":"2405.17485","n_code_links":0,"syntology":null},{"paper":"/paper/a-structure-aware-framework-for-learning","slug":"a-structure-aware-framework-for-learning","title":"A Structure-Aware Framework for Learning Device Placements on Computation Graphs","date":"2024-05-23","arxiv_id":"2405.14185","n_code_links":1,"syntology":null},{"paper":"/paper/ceebert-cross-domain-inference-in-early-exit","slug":"ceebert-cross-domain-inference-in-early-exit","title":"CEEBERT: Cross-Domain Inference in Early Exit BERT","date":"2024-05-23","arxiv_id":"2405.15039","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Div290/CeeBERT"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/hipporag-neurobiologically-inspired-long-term","slug":"hipporag-neurobiologically-inspired-long-term","title":"HippoRAG: Neurobiologically Inspired Long-Term Memory for Large Language Models","date":"2024-05-23","arxiv_id":"2405.14831","n_code_links":2,"syntology":{"ran":5,"of":5,"n_ran_checked":5,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["osu-nlp-group/hipporag"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"rafe-ranking-feedback-improves-query","title":"RaFe: Ranking Feedback Improves Query Rewriting for RAG","date":"2024-05-23","arxiv_id":"2405.14431","n_code_links":0,"syntology":null},{"paper":"/paper/vihatet5-enhancing-hate-speech-detection-in","slug":"vihatet5-enhancing-hate-speech-detection-in","title":"ViHateT5: Enhancing Hate Speech Detection in Vietnamese With A Unified Text-to-Text Transformer Model","date":"2024-05-23","arxiv_id":"2405.14141","n_code_links":1,"syntology":null},{"paper":"/paper/automated-evaluation-of-retrieval-augmented","slug":"automated-evaluation-of-retrieval-augmented","title":"Automated Evaluation of Retrieval-Augmented Language Models with Task-Specific Exam Generation","date":"2024-05-22","arxiv_id":"2405.13622","n_code_links":1,"syntology":null},{"paper":"/paper/flashrag-a-modular-toolkit-for-efficient","slug":"flashrag-a-modular-toolkit-for-efficient","title":"FlashRAG: A Modular Toolkit for Efficient Retrieval-Augmented Generation Research","date":"2024-05-22","arxiv_id":"2405.13576","n_code_links":1,"syntology":null},{"paper":"/paper/trojanrag-retrieval-augmented-generation-can","slug":"trojanrag-retrieval-augmented-generation-can","title":"TrojanRAG: Retrieval-Augmented Generation Can Be Backdoor Driver in Large Language Models","date":"2024-05-22","arxiv_id":"2405.13401","n_code_links":1,"syntology":null},{"paper":null,"slug":"unleashing-the-power-of-unlabeled-data-a-self","title":"Unleashing the Power of Unlabeled Data: A Self-supervised Learning Framework for Cyber Attack Detection in Smart Grids","date":"2024-05-22","arxiv_id":"2405.13965","n_code_links":0,"syntology":null},{"paper":null,"slug":"generative-ai-and-large-language-models-for","title":"Generative AI in Cybersecurity: A Comprehensive Review of LLM Applications and Vulnerabilities","date":"2024-05-21","arxiv_id":"2405.12750","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-reliable-ai-chatbots-are-for-disease","title":"How Reliable AI Chatbots are for Disease Prediction from Patient Complaints?","date":"2024-05-21","arxiv_id":"2405.13219","n_code_links":0,"syntology":null},{"paper":"/paper/the-2nd-futuredial-challenge-dialog-systems","slug":"the-2nd-futuredial-challenge-dialog-systems","title":"The 2nd FutureDial Challenge: Dialog Systems with Retrieval Augmented Generation (FutureDial-RAG)","date":"2024-05-21","arxiv_id":"2405.13084","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-review-on-the-use-of-large-language-models","title":"A review on the use of large language models as virtual tutors","date":"2024-05-20","arxiv_id":"2405.11983","n_code_links":0,"syntology":null},{"paper":null,"slug":"crema-crisis-response-through-computational","title":"CReMa: Crisis Response through Computational Identification and Matching of Cross-Lingual Requests and Offers Shared on Social Media","date":"2024-05-20","arxiv_id":"2405.11897","n_code_links":0,"syntology":null},{"paper":null,"slug":"degree-of-irrationality-sentiment-and-implied","title":"Degree of Irrationality: Sentiment and Implied Volatility Surface","date":"2024-05-20","arxiv_id":"2405.11730","n_code_links":0,"syntology":null},{"paper":null,"slug":"question-based-retrieval-using-atomic-units","title":"Question-Based Retrieval using Atomic Units for Enterprise RAG","date":"2024-05-20","arxiv_id":"2405.12363","n_code_links":0,"syntology":null},{"paper":null,"slug":"exploring-speech-style-spaces-with-language","title":"Exploring speech style spaces with language models: Emotional TTS without emotion labels","date":"2024-05-18","arxiv_id":"2405.11413","n_code_links":0,"syntology":null},{"paper":null,"slug":"activellm-large-language-model-based-active","title":"ActiveLLM: Large Language Model-based Active Learning for Textual Few-Shot Scenarios","date":"2024-05-17","arxiv_id":"2405.10808","n_code_links":0,"syntology":null},{"paper":null,"slug":"empowering-prior-to-court-legal-analysis-a","title":"Empowering Prior to Court Legal Analysis: A Transparent and Accessible Dataset for Defensive Statement Classification and Interpretation","date":"2024-05-17","arxiv_id":"2405.10702","n_code_links":0,"syntology":null},{"paper":null,"slug":"neuroassist-enhancing-cognitive-computer","title":"NeuroAssist: Enhancing Cognitive-Computer Synergy with Adaptive AI and Advanced Neural Decoding for Efficient EEG Signal Classification","date":"2024-05-17","arxiv_id":"2406.01600","n_code_links":0,"syntology":null},{"paper":"/paper/tailoring-vaccine-messaging-with-common","slug":"tailoring-vaccine-messaging-with-common","title":"Tailoring Vaccine Messaging with Common-Ground Opinions","date":"2024-05-17","arxiv_id":"2405.10861","n_code_links":1,"syntology":null},{"paper":null,"slug":"bridging-the-gap-in-online-hate-speech","title":"Bridging the gap in online hate speech detection: a comparative analysis of BERT and traditional models for homophobic content identification on X/Twitter","date":"2024-05-15","arxiv_id":"2405.09221","n_code_links":0,"syntology":null},{"paper":null,"slug":"im-rag-multi-round-retrieval-augmented","title":"IM-RAG: Multi-Round Retrieval-Augmented Generation Through Learning Inner Monologues","date":"2024-05-15","arxiv_id":"2405.13021","n_code_links":0,"syntology":null},{"paper":null,"slug":"transfer-learning-in-pre-trained-large","title":"Transfer Learning in Pre-Trained Large Language Models for Malware Detection Based on System Calls","date":"2024-05-15","arxiv_id":"2405.09318","n_code_links":0,"syntology":null},{"paper":null,"slug":"control-token-with-dense-passage-retrieval","title":"Control Token with Dense Passage Retrieval","date":"2024-05-13","arxiv_id":"2405.13008","n_code_links":0,"syntology":null},{"paper":"/paper/evaluation-of-retrieval-augmented-generation","slug":"evaluation-of-retrieval-augmented-generation","title":"Evaluation of Retrieval-Augmented Generation: A Survey","date":"2024-05-13","arxiv_id":"2405.07437","n_code_links":1,"syntology":null},{"paper":null,"slug":"from-questions-to-insightful-answers-building","title":"From Questions to Insightful Answers: Building an Informed Chatbot for University Resources","date":"2024-05-13","arxiv_id":"2405.08120","n_code_links":0,"syntology":null},{"paper":null,"slug":"duetrag-collaborative-retrieval-augmented","title":"DuetRAG: Collaborative Retrieval-Augmented Generation","date":"2024-05-12","arxiv_id":"2405.13002","n_code_links":0,"syntology":null},{"paper":null,"slug":"explainabledetector-exploring-transformer","title":"ExplainableDetector: Exploring Transformer-based Language Modeling Approach for SMS Spam Detection with Explainability Analysis","date":"2024-05-12","arxiv_id":"2405.08026","n_code_links":0,"syntology":null},{"paper":null,"slug":"l-u-pin-llm-based-political-ideology","title":"L(u)PIN: LLM-based Political Ideology Nowcasting","date":"2024-05-12","arxiv_id":"2405.07320","n_code_links":0,"syntology":null},{"paper":null,"slug":"tacoere-cluster-aware-compression-for-event","title":"TacoERE: Cluster-aware Compression for Event Relation Extraction","date":"2024-05-11","arxiv_id":"2405.06890","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-survey-on-rag-meets-llms-towards-retrieval","title":"A Survey on RAG Meeting LLMs: Towards Retrieval-Augmented Large Language Models","date":"2024-05-10","arxiv_id":"2405.06211","n_code_links":0,"syntology":null},{"paper":null,"slug":"canal-cyber-activity-news-alerting-language","title":"CANAL -- Cyber Activity News Alerting Language Model: Empirical Approach vs. Expensive LLM","date":"2024-05-10","arxiv_id":"2405.06772","n_code_links":0,"syntology":null},{"paper":null,"slug":"characterizing-the-accuracy-efficiency-trade","title":"Characterizing the Accuracy -- Efficiency Trade-off of Low-rank Decomposition in Language Models","date":"2024-05-10","arxiv_id":"2405.06626","n_code_links":0,"syntology":null},{"paper":"/paper/ditto-quantization-aware-secure-inference-of","slug":"ditto-quantization-aware-secure-inference-of","title":"Ditto: Quantization-aware Secure Inference of Transformers upon MPC","date":"2024-05-09","arxiv_id":"2405.05525","n_code_links":1,"syntology":null},{"paper":null,"slug":"reddit-impacts-a-named-entity-recognition","title":"Reddit-Impacts: A Named Entity Recognition Dataset for Analyzing Clinical and Social Effects of Substance Use Derived from Social Media","date":"2024-05-09","arxiv_id":"2405.06145","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-students-open-ended-written","title":"Evaluating Students' Open-ended Written Responses with LLMs: Using the RAG Framework for GPT-3.5, GPT-4, Claude-3, and Mistral-Large","date":"2024-05-08","arxiv_id":"2405.05444","n_code_links":0,"syntology":null},{"paper":null,"slug":"utilizing-large-language-models-to-generate","title":"Utilizing Large Language Models to Generate Synthetic Data to Increase the Performance of BERT-Based Neural Networks","date":"2024-05-08","arxiv_id":"2405.06695","n_code_links":0,"syntology":null},{"paper":"/paper/enriched-bert-embeddings-for-scholarly","slug":"enriched-bert-embeddings-for-scholarly","title":"Enriched BERT Embeddings for Scholarly Publication Classification","date":"2024-05-07","arxiv_id":"2405.04136","n_code_links":1,"syntology":null},{"paper":null,"slug":"eratta-extreme-rag-for-table-to-answers-with","title":"ERATTA: Extreme RAG for Table To Answers with Large Language Models","date":"2024-05-07","arxiv_id":"2405.03963","n_code_links":0,"syntology":null},{"paper":null,"slug":"evaluating-text-summaries-generated-by-large","title":"Evaluating Text Summaries Generated by Large Language Models Using OpenAI's GPT","date":"2024-05-07","arxiv_id":"2405.04053","n_code_links":0,"syntology":null},{"paper":null,"slug":"remote-diffusion","title":"Remote Diffusion","date":"2024-05-07","arxiv_id":"2405.04717","n_code_links":0,"syntology":null},{"paper":"/paper/revisiting-character-level-adversarial","slug":"revisiting-character-level-adversarial","title":"Revisiting Character-level Adversarial Attacks for Language Models","date":"2024-05-07","arxiv_id":"2405.04346","n_code_links":1,"syntology":{"ran":23,"of":31,"n_ran_checked":13,"n_instrument":10,"unverified":8,"pointer_only":0,"phrase":"23 ran (of which 3 constructed an object rather than computing a result; 13 with no instrument failure: 0 honoured, 0 violated, 13 with no contract checked; 10 where Syntology's instrument failed) · 8 unverified","official":{"repos":["lions-epfl/charmer"],"state":"official (archive's flag): 23 ran","n_ran":23,"n_constructed":3,"n_ran_no_instrument_failure":13,"n_unverified":8,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"robust-implementation-of-retrieval-augmented","title":"Robust Implementation of Retrieval-Augmented Generation on Edge-based Computing-in-Memory Architectures","date":"2024-05-07","arxiv_id":"2405.04700","n_code_links":0,"syntology":null},{"paper":null,"slug":"utilizing-gpt-to-enhance-text-summarization-a","title":"Utilizing GPT to Enhance Text Summarization: A Strategy to Minimize Hallucinations","date":"2024-05-07","arxiv_id":"2405.04039","n_code_links":0,"syntology":null},{"paper":null,"slug":"characterizing-the-dilemma-of-performance-and","title":"Characterizing the Dilemma of Performance and Index Size in Billion-Scale Vector Search and Breaking It with Second-Tier Memory","date":"2024-05-06","arxiv_id":"2405.03267","n_code_links":0,"syntology":null},{"paper":null,"slug":"compressing-long-context-for-enhancing-rag","title":"Compressing Long Context for Enhancing RAG with AMR-based Concept Distillation","date":"2024-05-06","arxiv_id":"2405.03085","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-android-malware-from-neural","title":"Detecting Android Malware: From Neural Embeddings to Hands-On Validation with BERTroid","date":"2024-05-06","arxiv_id":"2405.03620","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-anti-semitic-hate-speech-using","title":"Detecting Anti-Semitic Hate Speech using Transformer-based Large Language Models","date":"2024-05-06","arxiv_id":"2405.03794","n_code_links":0,"syntology":null},{"paper":"/paper/eragent-enhancing-retrieval-augmented","slug":"eragent-enhancing-retrieval-augmented","title":"ERAGent: Enhancing Retrieval-Augmented Language Models with Improved Accuracy, Efficiency, and Personalization","date":"2024-05-06","arxiv_id":"2405.06683","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-lecture-content-for-improved","title":"Leveraging Lecture Content for Improved Feedback: Explorations with GPT-4 and Retrieval Augmented Generation","date":"2024-05-05","arxiv_id":"2405.06681","n_code_links":0,"syntology":null},{"paper":null,"slug":"stochastic-rag-end-to-end-retrieval-augmented","title":"Stochastic RAG: End-to-End Retrieval-Augmented Generation through Expected Utility Maximization","date":"2024-05-05","arxiv_id":"2405.02816","n_code_links":0,"syntology":null},{"paper":"/paper/unraveling-the-dominance-of-large-language","slug":"unraveling-the-dominance-of-large-language","title":"Unraveling the Dominance of Large Language Models Over Transformer Models for Bangla Natural Language Inference: A Comprehensive Study","date":"2024-05-05","arxiv_id":"2405.02937","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-combination-of-bert-and-transformer-for","title":"A Combination of BERT and Transformer for Vietnamese Spelling Correction","date":"2024-05-04","arxiv_id":"2405.02573","n_code_links":0,"syntology":null},{"paper":null,"slug":"analyzing-narrative-processing-in-large","title":"Analyzing Narrative Processing in Large Language Models (LLMs): Using GPT4 to test BERT","date":"2024-05-03","arxiv_id":"2405.02024","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-retrieval-systems-in","title":"Comparative Analysis of Retrieval Systems in the Real World","date":"2024-05-03","arxiv_id":"2405.02048","n_code_links":0,"syntology":null},{"paper":"/paper/dallmi-domain-adaption-for-llm-based-multi","slug":"dallmi-domain-adaption-for-llm-based-multi","title":"DALLMi: Domain Adaption for LLM-based Multi-label Classifier","date":"2024-05-03","arxiv_id":"2405.01883","n_code_links":1,"syntology":null}],"record_sha256":"b53e7810f607dbf28d40ebed1f8550908339aa6ee5b82e3be430662c16b20154","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}