{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention/papers/101","list_of":"/method/attention","method":"Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":101,"pages_in_order":316,"rows_per_page":100,"rows":[10001,10100],"of":31583,"counts":{"archive_papers_tagged":31583,"with_a_code_link":13473,"where_syntology_ran_a_sample":3998,"not_listed_spam_title":0,"listed":31583,"listed_where_code_ran":3998,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3366,"every_run_a_failure_of_syntologys_instrument":632,"listed_with_a_run_with_no_instrument_failure":3366,"listed_every_run_a_failure_of_syntologys_instrument":632,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention","prev":"/method/attention/papers/100","next":"/method/attention/papers/102","papers":[{"paper":"/paper/rethinking-the-atmospheric-scattering-driven","slug":"rethinking-the-atmospheric-scattering-driven","title":"Rethinking the Atmospheric Scattering-driven Attention via Channel and Gamma Correction Priors for Low-Light Image Enhancement","date":"2024-09-09","arxiv_id":"2409.05274","n_code_links":1,"syntology":null},{"paper":"/paper/retrofitting-temporal-graph-neural-networks","slug":"retrofitting-temporal-graph-neural-networks","title":"Retrofitting Temporal Graph Neural Networks with Transformer","date":"2024-09-09","arxiv_id":"2409.05477","n_code_links":1,"syntology":null},{"paper":"/paper/revisiting-the-solution-of-meta-kdd-cup-2024","slug":"revisiting-the-solution-of-meta-kdd-cup-2024","title":"Revisiting the Solution of Meta KDD Cup 2024: CRAG","date":"2024-09-09","arxiv_id":"2409.15337","n_code_links":1,"syntology":null},{"paper":null,"slug":"rexuninlu-recursive-method-with-explicit","title":"RexUniNLU: Recursive Method with Explicit Schema Instructor for Universal NLU","date":"2024-09-09","arxiv_id":"2409.05275","n_code_links":0,"syntology":null},{"paper":"/paper/rotcatt-transunet-novel-deep-neural-network","slug":"rotcatt-transunet-novel-deep-neural-network","title":"RotCAtt-TransUNet++: Novel Deep Neural Network for Sophisticated Cardiac Segmentation","date":"2024-09-09","arxiv_id":"2409.05280","n_code_links":1,"syntology":null},{"paper":null,"slug":"sequential-posterior-sampling-with-diffusion","title":"Sequential Posterior Sampling with Diffusion Models","date":"2024-09-09","arxiv_id":"2409.05399","n_code_links":0,"syntology":null},{"paper":null,"slug":"songcreator-lyrics-based-universal-song","title":"SongCreator: Lyrics-based Universal Song Generation","date":"2024-09-09","arxiv_id":"2409.06029","n_code_links":0,"syntology":null},{"paper":null,"slug":"statistical-mechanics-of-min-max-problems","title":"Statistical Mechanics of Min-Max Problems","date":"2024-09-09","arxiv_id":"2409.06053","n_code_links":0,"syntology":null},{"paper":null,"slug":"sx-stitch-an-efficient-vms-unet-based","title":"SX-Stitch: An Efficient VMS-UNet Based Framework for Intraoperative Scoliosis X-Ray Image Stitching","date":"2024-09-09","arxiv_id":"2409.05681","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-building-a-robust-knowledge-intensive","title":"Towards Building a Robust Knowledge Intensive Question Answering Model with Large Language Models","date":"2024-09-09","arxiv_id":"2409.05385","n_code_links":0,"syntology":null},{"paper":null,"slug":"unveiling-induction-heads-provable-training","title":"Unveiling Induction Heads: Provable Training Dynamics and Feature Learning in Transformers","date":"2024-09-09","arxiv_id":"2409.10559","n_code_links":0,"syntology":null},{"paper":null,"slug":"zero-shot-outlier-detection-via-prior-data","title":"Zero-shot Outlier Detection via Prior-data Fitted Networks: Model Selection Bygone!","date":"2024-09-09","arxiv_id":"2409.05672","n_code_links":0,"syntology":null},{"paper":"/paper/a-survey-on-mixup-augmentations-and-beyond","slug":"a-survey-on-mixup-augmentations-and-beyond","title":"A Survey on Mixup Augmentations and Beyond","date":"2024-09-08","arxiv_id":"2409.05202","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-analog-and-digital-hybrid-attention","title":"An Analog and Digital Hybrid Attention Accelerator for Transformers with Charge-based In-memory Computing","date":"2024-09-08","arxiv_id":"2409.04940","n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-based-efficient-breath-sound","title":"Attention-Based Efficient Breath Sound Removal in Studio Audio Recordings","date":"2024-09-08","arxiv_id":"2409.04949","n_code_links":0,"syntology":null},{"paper":null,"slug":"audio-guided-fusion-techniques-for-multimodal","title":"Audio-Guided Fusion Techniques for Multimodal Emotion Analysis","date":"2024-09-08","arxiv_id":"2409.05007","n_code_links":0,"syntology":null},{"paper":null,"slug":"better-spanish-emotion-recognition-in-the","title":"Better Spanish Emotion Recognition In-the-wild: Bringing Attention to Deep Spectrum Voice Analysis","date":"2024-09-08","arxiv_id":"2409.05148","n_code_links":0,"syntology":null},{"paper":"/paper/dual-convolutional-neural-network-with","slug":"dual-convolutional-neural-network-with","title":"Dual convolutional neural network with attention for image blind denoising","date":"2024-09-08","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"instinfer-in-storage-attention-offloading-for","title":"InstInfer: In-Storage Attention Offloading for Cost-Effective Long-Context LLM Inference","date":"2024-09-08","arxiv_id":"2409.04992","n_code_links":0,"syntology":null},{"paper":null,"slug":"lung-detr-deformable-detection-transformer","title":"Lung-DETR: Deformable Detection Transformer for Sparse Lung Nodule Anomaly Detection","date":"2024-09-08","arxiv_id":"2409.05200","n_code_links":0,"syntology":null},{"paper":null,"slug":"mhs-stma-multimodal-hate-speech-detection-via","title":"MHS-STMA: Multimodal Hate Speech Detection via Scalable Transformer-Based Multilevel Attention Framework","date":"2024-09-08","arxiv_id":"2409.05136","n_code_links":0,"syntology":null},{"paper":"/paper/multi-v2x-a-large-scale-multi-modal-multi","slug":"multi-v2x-a-large-scale-multi-modal-multi","title":"Multi-V2X: A Large Scale Multi-modal Multi-penetration-rate Dataset for Cooperative Perception","date":"2024-09-08","arxiv_id":"2409.04980","n_code_links":1,"syntology":null},{"paper":"/paper/natias-neuron-attribution-based-transferable","slug":"natias-neuron-attribution-based-transferable","title":"Natias: Neuron Attribution based Transferable Image Adversarial Steganography","date":"2024-09-08","arxiv_id":"2409.04968","n_code_links":1,"syntology":null},{"paper":"/paper/onegen-efficient-one-pass-unified-generation","slug":"onegen-efficient-one-pass-unified-generation","title":"OneGen: Efficient One-Pass Unified Generation and Retrieval for LLMs","date":"2024-09-08","arxiv_id":"2409.05152","n_code_links":1,"syntology":null},{"paper":"/paper/pip-detecting-adversarial-examples-in-large","slug":"pip-detecting-adversarial-examples-in-large","title":"PIP: Detecting Adversarial Examples in Large Vision-Language Models via Attention Patterns of Irrelevant Probe Questions","date":"2024-09-08","arxiv_id":"2409.05076","n_code_links":1,"syntology":null},{"paper":"/paper/rcbevdet-toward-high-accuracy-radar-camera","slug":"rcbevdet-toward-high-accuracy-radar-camera","title":"RCBEVDet++: Toward High-accuracy Radar-Camera Fusion 3D Perception Network","date":"2024-09-08","arxiv_id":"2409.04979","n_code_links":0,"syntology":null},{"paper":"/paper/ss-brpe-self-supervised-blind-room-parameter","slug":"ss-brpe-self-supervised-blind-room-parameter","title":"SS-BRPE: Self-Supervised Blind Room Parameter Estimation Using Attention Mechanisms","date":"2024-09-08","arxiv_id":"2409.05212","n_code_links":1,"syntology":null},{"paper":"/paper/vision-fused-attack-advancing-aggressive-and","slug":"vision-fused-attack-advancing-aggressive-and","title":"Vision-fused Attack: Advancing Aggressive and Stealthy Adversarial Text against Neural Machine Translation","date":"2024-09-08","arxiv_id":"2409.05021","n_code_links":1,"syntology":{"ran":0,"of":4,"n_ran_checked":0,"n_instrument":0,"unverified":4,"pointer_only":4,"phrase":"0 ran · 4 unverified","official":{"repos":["levelower/vfa"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":4,"ran_from_kinds":[]}}},{"paper":null,"slug":"2409-13717","title":"DiVA-DocRE: A Discriminative and Voice-Aware Paradigm for Document-Level Relation Extraction","date":"2024-09-07","arxiv_id":"2409.13717","n_code_links":0,"syntology":null},{"paper":"/paper/activation-function-optimization-scheme-for","slug":"activation-function-optimization-scheme-for","title":"Activation Function Optimization Scheme for Image Classification","date":"2024-09-07","arxiv_id":"2409.04915","n_code_links":1,"syntology":null},{"paper":null,"slug":"adaptivefusion-adaptive-multi-modal-multi","title":"Towards Weather-Robust 3D Human Body Reconstruction: Millimeter-Wave Radar-Based Dataset, Benchmark, and Multi-Modal Fusion","date":"2024-09-07","arxiv_id":"2409.04851","n_code_links":0,"syntology":null},{"paper":null,"slug":"constrained-multi-layer-contrastive-learning","title":"Constrained Multi-Layer Contrastive Learning for Implicit Discourse Relationship Recognition","date":"2024-09-07","arxiv_id":"2409.13716","n_code_links":0,"syntology":null},{"paper":"/paper/cross-attention-inspired-selective-state","slug":"cross-attention-inspired-selective-state","title":"Cross-attention Inspired Selective State Space Models for Target Sound Extraction","date":"2024-09-07","arxiv_id":"2409.04803","n_code_links":1,"syntology":{"ran":3,"of":4,"n_ran_checked":2,"n_instrument":1,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["WuDH2000/CrossMamba"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"efficient-training-of-transformers-for","title":"Efficient Training of Transformers for Molecule Property Prediction on Small-scale Datasets","date":"2024-09-07","arxiv_id":"2409.04909","n_code_links":0,"syntology":null},{"paper":"/paper/fisheye-gs-lightweight-and-extensible","slug":"fisheye-gs-lightweight-and-extensible","title":"Fisheye-GS: Lightweight and Extensible Gaussian Splatting Module for Fisheye Cameras","date":"2024-09-07","arxiv_id":"2409.04751","n_code_links":0,"syntology":null},{"paper":"/paper/hullmi-human-vs-llm-identification-with","slug":"hullmi-human-vs-llm-identification-with","title":"HULLMI: Human vs LLM identification with explainability","date":"2024-09-07","arxiv_id":"2409.04808","n_code_links":1,"syntology":null},{"paper":null,"slug":"muap-multi-step-adaptive-prompt-learning-for","title":"MuAP: Multi-step Adaptive Prompt Learning for Vision-Language Model with Missing Modality","date":"2024-09-07","arxiv_id":"2409.04693","n_code_links":0,"syntology":null},{"paper":null,"slug":"naptune-efficient-model-tuning-for-mood","title":"NapTune: Efficient Model Tuning for Mood Classification using Previous Night's Sleep Measures along with Wearable Time-series","date":"2024-09-07","arxiv_id":"2409.04723","n_code_links":0,"syntology":null},{"paper":"/paper/sgseg-enabling-text-free-inference-in","slug":"sgseg-enabling-text-free-inference-in","title":"SGSeg: Enabling Text-free Inference in Language-guided Segmentation of Chest X-rays via Self-guidance","date":"2024-09-07","arxiv_id":"2409.04758","n_code_links":1,"syntology":null},{"paper":null,"slug":"spotactor-training-free-layout-controlled","title":"SpotActor: Training-Free Layout-Controlled Consistent Image Generation","date":"2024-09-07","arxiv_id":"2409.04801","n_code_links":0,"syntology":null},{"paper":null,"slug":"swin-transformer-for-robust-differentiation","title":"Swin Transformer for Robust Differentiation of Real and Synthetic Images: Intra- and Inter-Dataset Analysis","date":"2024-09-07","arxiv_id":"2409.04734","n_code_links":0,"syntology":null},{"paper":null,"slug":"top-gap-integrating-size-priors-in-cnns-for","title":"Top-GAP: Integrating Size Priors in CNNs for more Interpretability, Robustness, and Bias Mitigation","date":"2024-09-07","arxiv_id":"2409.04819","n_code_links":0,"syntology":null},{"paper":null,"slug":"training-free-style-consistent-image","title":"Training-Free Style Consistent Image Synthesis with Condition and Mask Guidance in E-Commerce","date":"2024-09-07","arxiv_id":"2409.04750","n_code_links":0,"syntology":null},{"paper":null,"slug":"unrolling-plug-and-play-network-for","title":"Unrolling Plug-and-Play Network for Hyperspectral Unmixing","date":"2024-09-07","arxiv_id":"2409.04719","n_code_links":0,"syntology":null},{"paper":null,"slug":"untie-the-knots-an-efficient-data","title":"Untie the Knots: An Efficient Data Augmentation Strategy for Long-Context Pre-Training in Language Models","date":"2024-09-07","arxiv_id":"2409.04774","n_code_links":0,"syntology":null},{"paper":null,"slug":"vidlpro-a-underline-vid-eo-underline-l","title":"VidLPRO: A $\\underline{Vid}$eo-$\\underline{L}$anguage $\\underline{P}$re-training Framework for $\\underline{Ro}$botic and Laparoscopic Surgery","date":"2024-09-07","arxiv_id":"2409.04732","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-first-look-at-efficient-and-secure-on","title":"A First Look At Efficient And Secure On-Device LLM Inference Against KV Leakage","date":"2024-09-06","arxiv_id":"2409.04040","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-novel-dataset-for-video-based-autism","title":"A Novel Dataset for Video-Based Autism Classification Leveraging Extra-Stimulatory Behavior","date":"2024-09-06","arxiv_id":"2409.04598","n_code_links":0,"syntology":null},{"paper":null,"slug":"actionflow-equivariant-accurate-and-efficient","title":"ActionFlow: Equivariant, Accurate, and Efficient Policies with Spatially Symmetric Flow Matching","date":"2024-09-06","arxiv_id":"2409.04576","n_code_links":0,"syntology":null},{"paper":null,"slug":"advancing-sem-based-nano-scale-defect","title":"Advancing SEM Based Nano-Scale Defect Analysis in Semiconductor Manufacturing for Advanced IC Nodes","date":"2024-09-06","arxiv_id":"2409.04310","n_code_links":0,"syntology":null},{"paper":"/paper/anymatch-efficient-zero-shot-entity-matching","slug":"anymatch-efficient-zero-shot-entity-matching","title":"AnyMatch -- Efficient Zero-Shot Entity Matching with a Small Language Model","date":"2024-09-06","arxiv_id":"2409.04073","n_code_links":1,"syntology":null},{"paper":null,"slug":"attentionx-exploiting-consensus-discrepancy","title":"AttentionX: Exploiting Consensus Discrepancy In Attention from A Distributed Optimization Perspective","date":"2024-09-06","arxiv_id":"2409.04275","n_code_links":0,"syntology":null},{"paper":null,"slug":"column-vocabulary-association-cva-semantic","title":"Column Vocabulary Association (CVA): semantic interpretation of dataless tables","date":"2024-09-06","arxiv_id":"2409.13709","n_code_links":0,"syntology":null},{"paper":null,"slug":"combining-llms-and-knowledge-graphs-to-reduce","title":"Combining LLMs and Knowledge Graphs to Reduce Hallucinations in Question Answering","date":"2024-09-06","arxiv_id":"2409.04181","n_code_links":0,"syntology":null},{"paper":"/paper/connectivity-inspired-network-for-context","slug":"connectivity-inspired-network-for-context","title":"Connectivity-Inspired Network for Context-Aware Recognition","date":"2024-09-06","arxiv_id":"2409.04360","n_code_links":1,"syntology":null},{"paper":null,"slug":"convolutional-transformer-based-image","title":"Convolutional Transformer-Based Image Compression","date":"2024-09-06","arxiv_id":"2409.04118","n_code_links":0,"syntology":null},{"paper":"/paper/dreamforge-motion-aware-autoregressive-video","slug":"dreamforge-motion-aware-autoregressive-video","title":"DreamForge: Motion-Aware Autoregressive Video Generation for Multi-View Driving Scenes","date":"2024-09-06","arxiv_id":"2409.04003","n_code_links":1,"syntology":null},{"paper":null,"slug":"galla-graph-aligned-large-language-models-for","title":"GALLa: Graph Aligned Large Language Models for Improved Source Code Understanding","date":"2024-09-06","arxiv_id":"2409.04183","n_code_links":0,"syntology":null},{"paper":null,"slug":"hermes-memory-efficient-pipeline-inference","title":"Hermes: Memory-Efficient Pipeline Inference for Large Models on Edge Devices","date":"2024-09-06","arxiv_id":"2409.04249","n_code_links":0,"syntology":null},{"paper":null,"slug":"protein-sequence-classification-using-natural","title":"Protein sequence classification using natural language processing techniques","date":"2024-09-06","arxiv_id":"2409.04491","n_code_links":0,"syntology":null},{"paper":"/paper/qihoo-t2x-an-efficiency-focused-diffusion","slug":"qihoo-t2x-an-efficiency-focused-diffusion","title":"Qihoo-T2X: An Efficient Proxy-Tokenized Diffusion Transformer for Text-to-Any-Task","date":"2024-09-06","arxiv_id":"2409.04005","n_code_links":1,"syntology":null},{"paper":null,"slug":"quantum-kernel-methods-under-scrutiny-a","title":"Quantum Kernel Methods under Scrutiny: A Benchmarking Study","date":"2024-09-06","arxiv_id":"2409.04406","n_code_links":0,"syntology":null},{"paper":null,"slug":"retrieval-augmented-generation-based-incident","title":"Retrieval Augmented Generation-Based Incident Resolution Recommendation System for IT Support","date":"2024-09-06","arxiv_id":"2409.13707","n_code_links":0,"syntology":null},{"paper":null,"slug":"searching-for-effective-preprocessing-method","title":"Searching for Effective Preprocessing Method and CNN-based Architecture with Efficient Channel Attention on Speech Emotion Recognition","date":"2024-09-06","arxiv_id":"2409.04007","n_code_links":0,"syntology":null},{"paper":null,"slug":"secure-traffic-sign-recognition-an-attention","title":"Secure Traffic Sign Recognition: An Attention-Enabled Universal Image Inpainting Mechanism against Light Patch Attacks","date":"2024-09-06","arxiv_id":"2409.04133","n_code_links":0,"syntology":null},{"paper":"/paper/specific-nucleic-acid-detection-using-a","slug":"specific-nucleic-acid-detection-using-a","title":"Specific Nucleic Acid Detection Using a Nanoparticle Hybridization Assay","date":"2024-09-06","arxiv_id":"2409.03983","n_code_links":2,"syntology":null},{"paper":"/paper/theory-analysis-and-best-practices-for","slug":"theory-analysis-and-best-practices-for","title":"Theory, Analysis, and Best Practices for Sigmoid Self-Attention","date":"2024-09-06","arxiv_id":"2409.04431","n_code_links":1,"syntology":{"ran":1,"of":2,"n_ran_checked":0,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["apple/ml-sigmoid-attention"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"towards-safer-online-spaces-simulating-and","title":"Towards Safer Online Spaces: Simulating and Assessing Intervention Strategies for Eating Disorder Discussions","date":"2024-09-06","arxiv_id":"2409.04043","n_code_links":0,"syntology":null},{"paper":null,"slug":"ui-jepa-towards-active-perception-of-user","title":"UI-JEPA: Towards Active Perception of User Intent through Onscreen User Activity","date":"2024-09-06","arxiv_id":"2409.04081","n_code_links":0,"syntology":null},{"paper":"/paper/unidet3d-multi-dataset-indoor-3d-object","slug":"unidet3d-multi-dataset-indoor-3d-object","title":"UniDet3D: Multi-dataset Indoor 3D Object Detection","date":"2024-09-06","arxiv_id":"2409.04234","n_code_links":1,"syntology":null},{"paper":null,"slug":"warpadam-a-new-adam-optimizer-based-on-meta","title":"WarpAdam: A new Adam optimizer based on Meta-Learning approach","date":"2024-09-06","arxiv_id":"2409.04244","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-different-level-text-protection-mechanism","title":"A Different Level Text Protection Mechanism With Differential Privacy","date":"2024-09-05","arxiv_id":"2409.03707","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multi-scale-analysis-of-the-czra","title":"A multi-scale analysis of the CzrA transcription repressor highlights the allosteric changes induced by metal ion binding","date":"2024-09-05","arxiv_id":"2409.03584","n_code_links":0,"syntology":null},{"paper":null,"slug":"active-fake-deepfake-camouflage","title":"Active Fake: DeepFake Camouflage","date":"2024-09-05","arxiv_id":"2409.03200","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-efficient-recommendation-model-based-on","title":"An Efficient Recommendation Model Based on Knowledge Graph Attention-Assisted Network (KGATAX)","date":"2024-09-05","arxiv_id":"2409.15315","n_code_links":0,"syntology":null},{"paper":"/paper/attend-first-consolidate-later-on-the","slug":"attend-first-consolidate-later-on-the","title":"Attend First, Consolidate Later: On the Importance of Attention in Different LLM Layers","date":"2024-09-05","arxiv_id":"2409.03621","n_code_links":1,"syntology":{"ran":6,"of":15,"n_ran_checked":6,"n_instrument":0,"unverified":9,"pointer_only":15,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 9 unverified","official":{"repos":["schwartz-lab-NLP/Attend-First-Consolidate-Later"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":9,"ran_from_kinds":["official"]}}},{"paper":"/paper/attention-heads-of-large-language-models-a","slug":"attention-heads-of-large-language-models-a","title":"Attention Heads of Large Language Models: A Survey","date":"2024-09-05","arxiv_id":"2409.03752","n_code_links":1,"syntology":null},{"paper":null,"slug":"bypassing-darcy-defense-indistinguishable","title":"Bypassing DARCY Defense: Indistinguishable Universal Adversarial Triggers","date":"2024-09-05","arxiv_id":"2409.03183","n_code_links":0,"syntology":null},{"paper":null,"slug":"ca-bert-leveraging-context-awareness-for","title":"CA-BERT: Leveraging Context Awareness for Enhanced Multi-Turn Chat Interaction","date":"2024-09-05","arxiv_id":"2409.13701","n_code_links":0,"syntology":null},{"paper":"/paper/cacer-clinical-concept-annotations-for-cancer","slug":"cacer-clinical-concept-annotations-for-cancer","title":"CACER: Clinical Concept Annotations for Cancer Events and Relations","date":"2024-09-05","arxiv_id":"2409.03905","n_code_links":1,"syntology":null},{"paper":"/paper/causal-temporal-representation-learning-with","slug":"causal-temporal-representation-learning-with","title":"Causal Temporal Representation Learning with Nonstationary Sparse Transition","date":"2024-09-05","arxiv_id":"2409.03142","n_code_links":1,"syntology":{"ran":7,"of":9,"n_ran_checked":4,"n_instrument":3,"unverified":2,"pointer_only":9,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","official":{"repos":["xiangchensong/ctrlns"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/characterizing-massive-activations-of","slug":"characterizing-massive-activations-of","title":"Characterizing Massive Activations of Attention Mechanism in Graph Neural Networks","date":"2024-09-05","arxiv_id":"2409.03463","n_code_links":1,"syntology":null},{"paper":"/paper/discovering-cyclists-street-visual","slug":"discovering-cyclists-street-visual","title":"Discovering Cyclists' Visual Preferences Through Shared Bike Trajectories and Street View Images Using Inverse Reinforcement Learning","date":"2024-09-05","arxiv_id":"2409.03148","n_code_links":1,"syntology":null},{"paper":"/paper/evaluating-open-source-sparse-autoencoders-on","slug":"evaluating-open-source-sparse-autoencoders-on","title":"Evaluating Open-Source Sparse Autoencoders on Disentangling Factual Knowledge in GPT-2 Small","date":"2024-09-05","arxiv_id":"2409.04478","n_code_links":1,"syntology":{"ran":1,"of":3,"n_ran_checked":0,"n_instrument":1,"unverified":2,"pointer_only":3,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 2 unverified","official":{"repos":["maheepchaudhary/sae-ravel"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"few-shot-continual-learning-for-activity","title":"Few-Shot Continual Learning for Activity Recognition in Classroom Surveillance Images","date":"2024-09-05","arxiv_id":"2409.03354","n_code_links":0,"syntology":null},{"paper":"/paper/hgamn-heterogeneous-graph-attention-matching","slug":"hgamn-heterogeneous-graph-attention-matching","title":"HGAMN: Heterogeneous Graph Attention Matching Network for Multilingual POI Retrieval at Baidu Maps","date":"2024-09-05","arxiv_id":"2409.03504","n_code_links":1,"syntology":null},{"paper":"/paper/lmlt-low-to-high-multi-level-vision","slug":"lmlt-low-to-high-multi-level-vision","title":"LMLT: Low-to-high Multi-Level Vision Transformer for Image Super-Resolution","date":"2024-09-05","arxiv_id":"2409.03516","n_code_links":1,"syntology":null},{"paper":null,"slug":"marags-a-multi-adapter-system-for-multi-task","title":"MARAGS: A Multi-Adapter System for Multi-Task Retrieval Augmented Generation Question Answering","date":"2024-09-05","arxiv_id":"2409.03171","n_code_links":0,"syntology":null},{"paper":null,"slug":"materialbench-evaluating-college-level","title":"MaterialBENCH: Evaluating College-Level Materials Science Problem-Solving Abilities of Large Language Models","date":"2024-09-05","arxiv_id":"2409.03161","n_code_links":0,"syntology":null},{"paper":"/paper/mvtn-a-multiscale-video-transformer-network","slug":"mvtn-a-multiscale-video-transformer-network","title":"MVTN: A Multiscale Video Transformer Network for Hand Gesture Recognition","date":"2024-09-05","arxiv_id":"2409.03890","n_code_links":1,"syntology":null},{"paper":"/paper/on-board-satellite-image-classification-for","slug":"on-board-satellite-image-classification-for","title":"Onboard Satellite Image Classification for Earth Observation: A Comparative Study of ViT Models","date":"2024-09-05","arxiv_id":"2409.03901","n_code_links":1,"syntology":null},{"paper":null,"slug":"rag-based-question-answering-for-contextual","title":"RAG based Question-Answering for Contextual Response Prediction System","date":"2024-09-05","arxiv_id":"2409.03708","n_code_links":0,"syntology":null},{"paper":null,"slug":"resultant-incremental-effectiveness-on","title":"Resultant: Incremental Effectiveness on Likelihood for Unsupervised Out-of-Distribution Detection","date":"2024-09-05","arxiv_id":"2409.03801","n_code_links":0,"syntology":null},{"paper":"/paper/revolutionizing-database-q-a-with-large","slug":"revolutionizing-database-q-a-with-large","title":"Revolutionizing Database Q&A with Large Language Models: Comprehensive Benchmark and Evaluation","date":"2024-09-05","arxiv_id":"2409.04475","n_code_links":1,"syntology":null},{"paper":null,"slug":"sketch-a-toolkit-for-streamlining-llm","title":"Sketch: A Toolkit for Streamlining LLM Operations","date":"2024-09-05","arxiv_id":"2409.03346","n_code_links":0,"syntology":null},{"paper":null,"slug":"spinmultinet-neural-network-potential","title":"SpinMultiNet: Neural Network Potential Incorporating Spin Degrees of Freedom with Multi-Task Learning","date":"2024-09-05","arxiv_id":"2409.03253","n_code_links":0,"syntology":null},{"paper":"/paper/surface-centric-modeling-for-high-fidelity","slug":"surface-centric-modeling-for-high-fidelity","title":"Surface-Centric Modeling for High-Fidelity Generalizable Neural Surface Reconstruction","date":"2024-09-05","arxiv_id":"2409.03634","n_code_links":1,"syntology":{"ran":5,"of":8,"n_ran_checked":5,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["prstrive/surf"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"tc-llava-rethinking-the-transfer-from-image","title":"TC-LLaVA: Rethinking the Transfer from Image to Video Understanding with Temporal Considerations","date":"2024-09-05","arxiv_id":"2409.03206","n_code_links":0,"syntology":null},{"paper":null,"slug":"why-mamba-is-effective-exploit-linear","title":"Why mamba is effective? Exploit Linear Transformer-Mamba Network for Multi-Modality Image Fusion","date":"2024-09-05","arxiv_id":"2409.03223","n_code_links":0,"syntology":null},{"paper":"/paper/xlam-a-family-of-large-action-models-to","slug":"xlam-a-family-of-large-action-models-to","title":"xLAM: A Family of Large Action Models to Empower AI Agent Systems","date":"2024-09-05","arxiv_id":"2409.03215","n_code_links":1,"syntology":null}],"record_sha256":"41d170cb445f547cc3e06daa8aede910cc98f429235ba790d82b4a446dc14206","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}