{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/multi-head-attention/papers/49","list_of":"/method/multi-head-attention","method":"Multi-Head Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":49,"pages_in_order":249,"rows_per_page":100,"rows":[4801,4900],"of":24855,"counts":{"archive_papers_tagged":24855,"with_a_code_link":11214,"where_syntology_ran_a_sample":3454,"not_listed_spam_title":0,"listed":24855,"listed_where_code_ran":3454,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2916,"every_run_a_failure_of_syntologys_instrument":538,"listed_with_a_run_with_no_instrument_failure":2916,"listed_every_run_a_failure_of_syntologys_instrument":538,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/multi-head-attention","prev":"/method/multi-head-attention/papers/48","next":"/method/multi-head-attention/papers/50","papers":[{"paper":null,"slug":"harmonic-reasoning-in-large-language-models","title":"Harmonic Reasoning in Large Language Models","date":"2024-09-09","arxiv_id":"2409.05521","n_code_links":0,"syntology":null},{"paper":null,"slug":"identifying-the-sources-of-ideological-bias","title":"Identifying the sources of ideological bias in GPT models through linguistic variation in output","date":"2024-09-09","arxiv_id":"2409.06043","n_code_links":0,"syntology":null},{"paper":"/paper/memorag-moving-towards-next-gen-rag-via","slug":"memorag-moving-towards-next-gen-rag-via","title":"MemoRAG: Moving towards Next-Gen RAG Via Memory-Inspired Knowledge Discovery","date":"2024-09-09","arxiv_id":"2409.05591","n_code_links":1,"syntology":{"ran":9,"of":11,"n_ran_checked":6,"n_instrument":3,"unverified":2,"pointer_only":6,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","official":{"repos":["qhjqhj00/memorag"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"qibert-classifying-online-conversations","title":"QiBERT -- Classifying Online Conversations Messages with BERT as a Feature","date":"2024-09-09","arxiv_id":"2409.05530","n_code_links":0,"syntology":null},{"paper":null,"slug":"regression-with-large-language-models-for","title":"Regression with Large Language Models for Materials and Molecular Property Prediction","date":"2024-09-09","arxiv_id":"2409.06080","n_code_links":0,"syntology":null},{"paper":"/paper/retrofitting-temporal-graph-neural-networks","slug":"retrofitting-temporal-graph-neural-networks","title":"Retrofitting Temporal Graph Neural Networks with Transformer","date":"2024-09-09","arxiv_id":"2409.05477","n_code_links":1,"syntology":null},{"paper":"/paper/revisiting-the-solution-of-meta-kdd-cup-2024","slug":"revisiting-the-solution-of-meta-kdd-cup-2024","title":"Revisiting the Solution of Meta KDD Cup 2024: CRAG","date":"2024-09-09","arxiv_id":"2409.15337","n_code_links":1,"syntology":null},{"paper":"/paper/rotcatt-transunet-novel-deep-neural-network","slug":"rotcatt-transunet-novel-deep-neural-network","title":"RotCAtt-TransUNet++: Novel Deep Neural Network for Sophisticated Cardiac Segmentation","date":"2024-09-09","arxiv_id":"2409.05280","n_code_links":1,"syntology":null},{"paper":null,"slug":"sequential-posterior-sampling-with-diffusion","title":"Sequential Posterior Sampling with Diffusion Models","date":"2024-09-09","arxiv_id":"2409.05399","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-building-a-robust-knowledge-intensive","title":"Towards Building a Robust Knowledge Intensive Question Answering Model with Large Language Models","date":"2024-09-09","arxiv_id":"2409.05385","n_code_links":0,"syntology":null},{"paper":null,"slug":"zero-shot-outlier-detection-via-prior-data","title":"Zero-shot Outlier Detection via Prior-data Fitted Networks: Model Selection Bygone!","date":"2024-09-09","arxiv_id":"2409.05672","n_code_links":0,"syntology":null},{"paper":null,"slug":"audio-guided-fusion-techniques-for-multimodal","title":"Audio-Guided Fusion Techniques for Multimodal Emotion Analysis","date":"2024-09-08","arxiv_id":"2409.05007","n_code_links":0,"syntology":null},{"paper":null,"slug":"lung-detr-deformable-detection-transformer","title":"Lung-DETR: Deformable Detection Transformer for Sparse Lung Nodule Anomaly Detection","date":"2024-09-08","arxiv_id":"2409.05200","n_code_links":0,"syntology":null},{"paper":"/paper/onegen-efficient-one-pass-unified-generation","slug":"onegen-efficient-one-pass-unified-generation","title":"OneGen: Efficient One-Pass Unified Generation and Retrieval for LLMs","date":"2024-09-08","arxiv_id":"2409.05152","n_code_links":1,"syntology":null},{"paper":"/paper/vision-fused-attack-advancing-aggressive-and","slug":"vision-fused-attack-advancing-aggressive-and","title":"Vision-fused Attack: Advancing Aggressive and Stealthy Adversarial Text against Neural Machine Translation","date":"2024-09-08","arxiv_id":"2409.05021","n_code_links":1,"syntology":{"ran":0,"of":4,"n_ran_checked":0,"n_instrument":0,"unverified":4,"pointer_only":4,"phrase":"0 ran · 4 unverified","official":{"repos":["levelower/vfa"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":4,"ran_from_kinds":[]}}},{"paper":"/paper/activation-function-optimization-scheme-for","slug":"activation-function-optimization-scheme-for","title":"Activation Function Optimization Scheme for Image Classification","date":"2024-09-07","arxiv_id":"2409.04915","n_code_links":1,"syntology":null},{"paper":null,"slug":"adaptivefusion-adaptive-multi-modal-multi","title":"Towards Weather-Robust 3D Human Body Reconstruction: Millimeter-Wave Radar-Based Dataset, Benchmark, and Multi-Modal Fusion","date":"2024-09-07","arxiv_id":"2409.04851","n_code_links":0,"syntology":null},{"paper":null,"slug":"constrained-multi-layer-contrastive-learning","title":"Constrained Multi-Layer Contrastive Learning for Implicit Discourse Relationship Recognition","date":"2024-09-07","arxiv_id":"2409.13716","n_code_links":0,"syntology":null},{"paper":"/paper/cross-attention-inspired-selective-state","slug":"cross-attention-inspired-selective-state","title":"Cross-attention Inspired Selective State Space Models for Target Sound Extraction","date":"2024-09-07","arxiv_id":"2409.04803","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":2,"n_instrument":1,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["WuDH2000/CrossMamba"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"efficient-training-of-transformers-for","title":"Efficient Training of Transformers for Molecule Property Prediction on Small-scale Datasets","date":"2024-09-07","arxiv_id":"2409.04909","n_code_links":0,"syntology":null},{"paper":null,"slug":"muap-multi-step-adaptive-prompt-learning-for","title":"MuAP: Multi-step Adaptive Prompt Learning for Vision-Language Model with Missing Modality","date":"2024-09-07","arxiv_id":"2409.04693","n_code_links":0,"syntology":null},{"paper":null,"slug":"naptune-efficient-model-tuning-for-mood","title":"NapTune: Efficient Model Tuning for Mood Classification using Previous Night's Sleep Measures along with Wearable Time-series","date":"2024-09-07","arxiv_id":"2409.04723","n_code_links":0,"syntology":null},{"paper":null,"slug":"swin-transformer-for-robust-differentiation","title":"Swin Transformer for Robust Differentiation of Real and Synthetic Images: Intra- and Inter-Dataset Analysis","date":"2024-09-07","arxiv_id":"2409.04734","n_code_links":0,"syntology":null},{"paper":null,"slug":"vidlpro-a-underline-vid-eo-underline-l","title":"VidLPRO: A $\\underline{Vid}$eo-$\\underline{L}$anguage $\\underline{P}$re-training Framework for $\\underline{Ro}$botic and Laparoscopic Surgery","date":"2024-09-07","arxiv_id":"2409.04732","n_code_links":0,"syntology":null},{"paper":null,"slug":"actionflow-equivariant-accurate-and-efficient","title":"ActionFlow: Equivariant, Accurate, and Efficient Policies with Spatially Symmetric Flow Matching","date":"2024-09-06","arxiv_id":"2409.04576","n_code_links":0,"syntology":null},{"paper":null,"slug":"advancing-sem-based-nano-scale-defect","title":"Advancing SEM Based Nano-Scale Defect Analysis in Semiconductor Manufacturing for Advanced IC Nodes","date":"2024-09-06","arxiv_id":"2409.04310","n_code_links":0,"syntology":null},{"paper":"/paper/anymatch-efficient-zero-shot-entity-matching","slug":"anymatch-efficient-zero-shot-entity-matching","title":"AnyMatch -- Efficient Zero-Shot Entity Matching with a Small Language Model","date":"2024-09-06","arxiv_id":"2409.04073","n_code_links":1,"syntology":null},{"paper":null,"slug":"column-vocabulary-association-cva-semantic","title":"Column Vocabulary Association (CVA): semantic interpretation of dataless tables","date":"2024-09-06","arxiv_id":"2409.13709","n_code_links":0,"syntology":null},{"paper":null,"slug":"combining-llms-and-knowledge-graphs-to-reduce","title":"Combining LLMs and Knowledge Graphs to Reduce Hallucinations in Question Answering","date":"2024-09-06","arxiv_id":"2409.04181","n_code_links":0,"syntology":null},{"paper":null,"slug":"convolutional-transformer-based-image","title":"Convolutional Transformer-Based Image Compression","date":"2024-09-06","arxiv_id":"2409.04118","n_code_links":0,"syntology":null},{"paper":null,"slug":"galla-graph-aligned-large-language-models-for","title":"GALLa: Graph Aligned Large Language Models for Improved Source Code Understanding","date":"2024-09-06","arxiv_id":"2409.04183","n_code_links":0,"syntology":null},{"paper":null,"slug":"hermes-memory-efficient-pipeline-inference","title":"Hermes: Memory-Efficient Pipeline Inference for Large Models on Edge Devices","date":"2024-09-06","arxiv_id":"2409.04249","n_code_links":0,"syntology":null},{"paper":null,"slug":"protein-sequence-classification-using-natural","title":"Protein sequence classification using natural language processing techniques","date":"2024-09-06","arxiv_id":"2409.04491","n_code_links":0,"syntology":null},{"paper":"/paper/qihoo-t2x-an-efficiency-focused-diffusion","slug":"qihoo-t2x-an-efficiency-focused-diffusion","title":"Qihoo-T2X: An Efficient Proxy-Tokenized Diffusion Transformer for Text-to-Any-Task","date":"2024-09-06","arxiv_id":"2409.04005","n_code_links":1,"syntology":null},{"paper":null,"slug":"retrieval-augmented-generation-based-incident","title":"Retrieval Augmented Generation-Based Incident Resolution Recommendation System for IT Support","date":"2024-09-06","arxiv_id":"2409.13707","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-safer-online-spaces-simulating-and","title":"Towards Safer Online Spaces: Simulating and Assessing Intervention Strategies for Eating Disorder Discussions","date":"2024-09-06","arxiv_id":"2409.04043","n_code_links":0,"syntology":null},{"paper":null,"slug":"ui-jepa-towards-active-perception-of-user","title":"UI-JEPA: Towards Active Perception of User Intent through Onscreen User Activity","date":"2024-09-06","arxiv_id":"2409.04081","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-different-level-text-protection-mechanism","title":"A Different Level Text Protection Mechanism With Differential Privacy","date":"2024-09-05","arxiv_id":"2409.03707","n_code_links":0,"syntology":null},{"paper":null,"slug":"bypassing-darcy-defense-indistinguishable","title":"Bypassing DARCY Defense: Indistinguishable Universal Adversarial Triggers","date":"2024-09-05","arxiv_id":"2409.03183","n_code_links":0,"syntology":null},{"paper":null,"slug":"ca-bert-leveraging-context-awareness-for","title":"CA-BERT: Leveraging Context Awareness for Enhanced Multi-Turn Chat Interaction","date":"2024-09-05","arxiv_id":"2409.13701","n_code_links":0,"syntology":null},{"paper":"/paper/cacer-clinical-concept-annotations-for-cancer","slug":"cacer-clinical-concept-annotations-for-cancer","title":"CACER: Clinical Concept Annotations for Cancer Events and Relations","date":"2024-09-05","arxiv_id":"2409.03905","n_code_links":1,"syntology":null},{"paper":"/paper/causal-temporal-representation-learning-with","slug":"causal-temporal-representation-learning-with","title":"Causal Temporal Representation Learning with Nonstationary Sparse Transition","date":"2024-09-05","arxiv_id":"2409.03142","n_code_links":1,"syntology":{"ran":7,"of":9,"n_ran_checked":4,"n_instrument":3,"unverified":2,"pointer_only":9,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","official":{"repos":["xiangchensong/ctrlns"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/characterizing-massive-activations-of","slug":"characterizing-massive-activations-of","title":"Characterizing Massive Activations of Attention Mechanism in Graph Neural Networks","date":"2024-09-05","arxiv_id":"2409.03463","n_code_links":1,"syntology":null},{"paper":"/paper/evaluating-open-source-sparse-autoencoders-on","slug":"evaluating-open-source-sparse-autoencoders-on","title":"Evaluating Open-Source Sparse Autoencoders on Disentangling Factual Knowledge in GPT-2 Small","date":"2024-09-05","arxiv_id":"2409.04478","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":0,"n_instrument":1,"unverified":2,"pointer_only":3,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["maheepchaudhary/sae-ravel"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/lmlt-low-to-high-multi-level-vision","slug":"lmlt-low-to-high-multi-level-vision","title":"LMLT: Low-to-high Multi-Level Vision Transformer for Image Super-Resolution","date":"2024-09-05","arxiv_id":"2409.03516","n_code_links":1,"syntology":null},{"paper":null,"slug":"marags-a-multi-adapter-system-for-multi-task","title":"MARAGS: A Multi-Adapter System for Multi-Task Retrieval Augmented Generation Question Answering","date":"2024-09-05","arxiv_id":"2409.03171","n_code_links":0,"syntology":null},{"paper":null,"slug":"materialbench-evaluating-college-level","title":"MaterialBENCH: Evaluating College-Level Materials Science Problem-Solving Abilities of Large Language Models","date":"2024-09-05","arxiv_id":"2409.03161","n_code_links":0,"syntology":null},{"paper":"/paper/mvtn-a-multiscale-video-transformer-network","slug":"mvtn-a-multiscale-video-transformer-network","title":"MVTN: A Multiscale Video Transformer Network for Hand Gesture Recognition","date":"2024-09-05","arxiv_id":"2409.03890","n_code_links":1,"syntology":null},{"paper":"/paper/on-board-satellite-image-classification-for","slug":"on-board-satellite-image-classification-for","title":"Onboard Satellite Image Classification for Earth Observation: A Comparative Study of ViT Models","date":"2024-09-05","arxiv_id":"2409.03901","n_code_links":1,"syntology":null},{"paper":null,"slug":"rag-based-question-answering-for-contextual","title":"RAG based Question-Answering for Contextual Response Prediction System","date":"2024-09-05","arxiv_id":"2409.03708","n_code_links":0,"syntology":null},{"paper":"/paper/revolutionizing-database-q-a-with-large","slug":"revolutionizing-database-q-a-with-large","title":"Revolutionizing Database Q&A with Large Language Models: Comprehensive Benchmark and Evaluation","date":"2024-09-05","arxiv_id":"2409.04475","n_code_links":1,"syntology":null},{"paper":null,"slug":"sketch-a-toolkit-for-streamlining-llm","title":"Sketch: A Toolkit for Streamlining LLM Operations","date":"2024-09-05","arxiv_id":"2409.03346","n_code_links":0,"syntology":null},{"paper":null,"slug":"why-mamba-is-effective-exploit-linear","title":"Why mamba is effective? Exploit Linear Transformer-Mamba Network for Multi-Modality Image Fusion","date":"2024-09-05","arxiv_id":"2409.03223","n_code_links":0,"syntology":null},{"paper":"/paper/xlam-a-family-of-large-action-models-to","slug":"xlam-a-family-of-large-action-models-to","title":"xLAM: A Family of Large Action Models to Empower AI Agent Systems","date":"2024-09-05","arxiv_id":"2409.03215","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-comparative-study-on-large-language-models","title":"A Comparative Study on Large Language Models for Log Parsing","date":"2024-09-04","arxiv_id":"2409.02474","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-data-selection-approach-for-enhancing-low","title":"A Data Selection Approach for Enhancing Low Resource Machine Translation Using Cross-Lingual Sentence Representations","date":"2024-09-04","arxiv_id":"2409.02712","n_code_links":0,"syntology":null},{"paper":null,"slug":"advancing-cyber-incident-timeline-analysis","title":"GenDFIR: Advancing Cyber Incident Timeline Analysis Through Retrieval Augmented Generation and Large Language Models","date":"2024-09-04","arxiv_id":"2409.02572","n_code_links":0,"syntology":null},{"paper":null,"slug":"causality-aware-transformer-networks-for","title":"Causality-Aware Transformer Networks for Robotic Navigation","date":"2024-09-04","arxiv_id":"2409.02669","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-calls-to-action-in-multimodal","title":"Detecting Calls to Action in Multimodal Content: Analysis of the 2021 German Federal Election Campaign on Instagram","date":"2024-09-04","arxiv_id":"2409.02690","n_code_links":0,"syntology":null},{"paper":null,"slug":"diversify-verify-adapt-efficient-and-robust","title":"Diversify-verify-adapt: Efficient and Robust Retrieval-Augmented Ambiguous Question Answering","date":"2024-09-04","arxiv_id":"2409.02361","n_code_links":0,"syntology":null},{"paper":null,"slug":"historical-german-text-normalization-using","title":"Historical German Text Normalization Using Type- and Token-Based Language Modeling","date":"2024-09-04","arxiv_id":"2409.02841","n_code_links":0,"syntology":null},{"paper":"/paper/how-dreams-are-made-emulating-satellite","slug":"how-dreams-are-made-emulating-satellite","title":"How DREAMS are made: Emulating Satellite Galaxy and Subhalo Populations with Diffusion Models and Point Clouds","date":"2024-09-04","arxiv_id":"2409.02980","n_code_links":1,"syntology":{"ran":10,"of":11,"n_ran_checked":10,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["trivnguyen/nehod_torch"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"how-privacy-savvy-are-large-language-models-a","title":"How Privacy-Savvy Are Large Language Models? A Case Study on Compliance and Privacy Technical Review","date":"2024-09-04","arxiv_id":"2409.02375","n_code_links":0,"syntology":null},{"paper":"/paper/hypothesizing-missing-causal-variables-with","slug":"hypothesizing-missing-causal-variables-with","title":"Hypothesizing Missing Causal Variables with LLMs","date":"2024-09-04","arxiv_id":"2409.02604","n_code_links":1,"syntology":null},{"paper":null,"slug":"iconformer-dynamic-parameter-efficient-tuning","title":"iConFormer: Dynamic Parameter-Efficient Tuning with Input-Conditioned Adaptation","date":"2024-09-04","arxiv_id":"2409.02838","n_code_links":0,"syntology":null},{"paper":null,"slug":"incorporating-like-minded-peers-to-overcome","title":"Incorporating Like-Minded Peers to Overcome Friend Data Sparsity in Session-Based Social Recommendations","date":"2024-09-04","arxiv_id":"2409.02702","n_code_links":0,"syntology":null},{"paper":null,"slug":"irrelevant-alternatives-bias-large-language","title":"Irrelevant Alternatives Bias Large Language Model Hiring Decisions","date":"2024-09-04","arxiv_id":"2409.15299","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-as-efficient-reward","title":"Large Language Models as Efficient Reward Function Searchers for Custom-Environment Multi-Objective Reinforcement Learning","date":"2024-09-04","arxiv_id":"2409.02428","n_code_links":0,"syntology":null},{"paper":null,"slug":"leveraging-interpretability-in-the","title":"Leveraging Interpretability in the Transformer to Automate the Proactive Scaling of Cloud Resources","date":"2024-09-04","arxiv_id":"2409.03103","n_code_links":0,"syntology":null},{"paper":"/paper/longllava-scaling-multi-modal-llms-to-1000","slug":"longllava-scaling-multi-modal-llms-to-1000","title":"LongLLaVA: Scaling Multi-modal LLMs to 1000 Images Efficiently via a Hybrid Architecture","date":"2024-09-04","arxiv_id":"2409.02889","n_code_links":1,"syntology":{"ran":7,"of":7,"n_ran_checked":3,"n_instrument":4,"unverified":0,"pointer_only":7,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 4 where Syntology's instrument failed) · 0 unverified","official":{"repos":["freedomintelligence/longllava"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"moa-is-all-you-need-building-llm-research","title":"MoA is All You Need: Building LLM Research Team using Mixture of Agents","date":"2024-09-04","arxiv_id":"2409.07487","n_code_links":0,"syntology":null},{"paper":"/paper/mobileunetr-a-lightweight-end-to-end-hybrid","slug":"mobileunetr-a-lightweight-end-to-end-hybrid","title":"MobileUNETR: A Lightweight End-To-End Hybrid Vision Transformer For Efficient Medical Image Segmentation","date":"2024-09-04","arxiv_id":"2409.03062","n_code_links":1,"syntology":null},{"paper":"/paper/more-is-more-addition-bias-in-large-language","slug":"more-is-more-addition-bias-in-large-language","title":"More is More: Addition Bias in Large Language Models","date":"2024-09-04","arxiv_id":"2409.02569","n_code_links":1,"syntology":null},{"paper":null,"slug":"mosmos-multi-organ-segmentation-facilitated","title":"MOSMOS: Multi-organ segmentation facilitated by medical report supervision","date":"2024-09-04","arxiv_id":"2409.02418","n_code_links":0,"syntology":null},{"paper":"/paper/multi-head-attention-residual-unfolded","slug":"multi-head-attention-residual-unfolded","title":"Multi-Head Attention Residual Unfolded Network for Model-Based Pansharpening","date":"2024-09-04","arxiv_id":"2409.02675","n_code_links":1,"syntology":null},{"paper":null,"slug":"openfact-at-checkthat-2024-combining-multiple","title":"OpenFact at CheckThat! 2024: Combining Multiple Attack Methods for Effective Adversarial Text Generation","date":"2024-09-04","arxiv_id":"2409.02649","n_code_links":0,"syntology":null},{"paper":null,"slug":"pre-training-data-selection-for-biomedical","title":"Pre-training data selection for biomedical domain adaptation using journal impact metrics","date":"2024-09-04","arxiv_id":"2409.02725","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-text-to-cypher-using-combination-of","title":"Robust Text-to-Cypher Using Combination of BERT, GraphSAGE, and Transformer (CoBGT) Model","date":"2024-09-04","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-data-centric-face-anti-spoofing","title":"Towards Data-Centric Face Anti-Spoofing: Improving Cross-domain Generalization via Physics-based Data Synthesis","date":"2024-09-04","arxiv_id":"2409.03501","n_code_links":0,"syntology":null},{"paper":null,"slug":"training-universal-vocoders-with-feature","title":"Training Universal Vocoders with Feature Smoothing-Based Augmentation Methods for High-Quality TTS Systems","date":"2024-09-04","arxiv_id":"2409.02517","n_code_links":0,"syntology":null},{"paper":null,"slug":"waveletgpt-wavelets-meet-large-language","title":"Wavelet GPT: Wavelet Inspired Large Language Models","date":"2024-09-04","arxiv_id":"2409.12924","n_code_links":0,"syntology":null},{"paper":null,"slug":"1dcnntrans-bisindo-sign-language-interpreters","title":"1DCNNTrans: BISINDO Sign Language Interpreters in Improving the Inclusiveness of Public Services","date":"2024-09-03","arxiv_id":"2409.01975","n_code_links":0,"syntology":null},{"paper":"/paper/2409-13694","slug":"2409-13694","title":"Multi-Source Knowledge Pruning for Retrieval-Augmented Generation: A Benchmark and Empirical Study","date":"2024-09-03","arxiv_id":"2409.13694","n_code_links":2,"syntology":null},{"paper":"/paper/2409-13695","slug":"2409-13695","title":"You Only Use Reactive Attention Slice For Long Context Retrieval","date":"2024-09-03","arxiv_id":"2409.13695","n_code_links":1,"syntology":null},{"paper":null,"slug":"adacomp-extractive-context-compression-with","title":"AdaComp: Extractive Context Compression with Adaptive Predictor for Retrieval-Augmented Large Language Models","date":"2024-09-03","arxiv_id":"2409.01579","n_code_links":0,"syntology":null},{"paper":null,"slug":"astromae-redshift-prediction-using-a-masked","title":"AstroMAE: Redshift Prediction Using a Masked Autoencoder with a Novel Fine-Tuning Architecture","date":"2024-09-03","arxiv_id":"2409.01825","n_code_links":0,"syntology":null},{"paper":null,"slug":"beaver-an-enterprise-benchmark-for-text-to","title":"BEAVER: An Enterprise Benchmark for Text-to-SQL","date":"2024-09-03","arxiv_id":"2409.02038","n_code_links":0,"syntology":null},{"paper":null,"slug":"benchmarking-cognitive-domains-for-llms","title":"Benchmarking Cognitive Domains for LLMs: Insights from Taiwanese Hakka Culture","date":"2024-09-03","arxiv_id":"2409.01556","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-learning-based","title":"Comparative Analysis of Learning-Based Methods for Transient Stability Assessment","date":"2024-09-03","arxiv_id":"2409.02336","n_code_links":0,"syntology":null},{"paper":null,"slug":"dialogue-you-can-trust-human-and-ai","title":"Dialogue You Can Trust: Human and AI Perspectives on Generated Conversations","date":"2024-09-03","arxiv_id":"2409.01808","n_code_links":0,"syntology":null},{"paper":null,"slug":"f2former-when-fractional-fourier-meets-deep","title":"F2former: When Fractional Fourier Meets Deep Wiener Deconvolution and Selective Frequency Transformer for Image Deblurring","date":"2024-09-03","arxiv_id":"2409.02056","n_code_links":0,"syntology":null},{"paper":"/paper/frequency-spatial-entanglement-learning-for","slug":"frequency-spatial-entanglement-learning-for","title":"Frequency-Spatial Entanglement Learning for Camouflaged Object Detection","date":"2024-09-03","arxiv_id":"2409.01686","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":1,"n_instrument":2,"unverified":3,"pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 3 unverified","official":{"repos":["csysi/fsel"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"improving-apple-object-detection-with","title":"Improving Apple Object Detection with Occlusion-Enhanced Distillation","date":"2024-09-03","arxiv_id":"2409.01573","n_code_links":0,"syntology":null},{"paper":null,"slug":"in-defense-of-rag-in-the-era-of-long-context","title":"In Defense of RAG in the Era of Long-Context Language Models","date":"2024-09-03","arxiv_id":"2409.01666","n_code_links":0,"syntology":null},{"paper":"/paper/it-is-time-to-develop-an-auditing-framework","slug":"it-is-time-to-develop-an-auditing-framework","title":"It is Time to Develop an Auditing Framework to Promote Value Aware Chatbots","date":"2024-09-03","arxiv_id":"2409.01539","n_code_links":1,"syntology":null},{"paper":null,"slug":"leveraging-large-language-models-for-solving","title":"Leveraging Large Language Models for Solving Rare MIP Challenges","date":"2024-09-03","arxiv_id":"2409.04464","n_code_links":0,"syntology":null},{"paper":"/paper/lifegpt-topology-agnostic-generative","slug":"lifegpt-topology-agnostic-generative","title":"LifeGPT: Topology-Agnostic Generative Pretrained Transformer Model for Cellular Automata","date":"2024-09-03","arxiv_id":"2409.12182","n_code_links":1,"syntology":null},{"paper":"/paper/multi-modal-adapter-for-vision-language","slug":"multi-modal-adapter-for-vision-language","title":"Multi-Modal Adapter for Vision-Language Models","date":"2024-09-03","arxiv_id":"2409.02958","n_code_links":1,"syntology":null},{"paper":null,"slug":"on-the-design-space-between-transformers-and","title":"On the Design Space Between Transformers and Recursive Neural Nets","date":"2024-09-03","arxiv_id":"2409.01531","n_code_links":0,"syntology":null},{"paper":null,"slug":"pmt-mae-dual-branch-self-supervised-learning","title":"PMT-MAE: Dual-Branch Self-Supervised Learning with Distillation for Efficient Point Cloud Classification","date":"2024-09-03","arxiv_id":"2409.02007","n_code_links":0,"syntology":null}],"record_sha256":"3d47296a712f3743442f4f475077ba077e3cff7ea922451a42ca484f240c7d23","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}