{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention/papers/72","list_of":"/method/attention","method":"Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":72,"pages_in_order":316,"rows_per_page":100,"rows":[7101,7200],"of":31583,"counts":{"archive_papers_tagged":31583,"with_a_code_link":13473,"where_syntology_ran_a_sample":3998,"not_listed_spam_title":0,"listed":31583,"listed_where_code_ran":3998,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3366,"every_run_a_failure_of_syntologys_instrument":632,"listed_with_a_run_with_no_instrument_failure":3366,"listed_every_run_a_failure_of_syntologys_instrument":632,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention","prev":"/method/attention/papers/71","next":"/method/attention/papers/73","papers":[{"paper":null,"slug":"unraveling-arithmetic-in-large-language","title":"Unraveling Arithmetic in Large Language Models: The Role of Algebraic Structures","date":"2024-11-25","arxiv_id":"2411.16260","n_code_links":0,"syntology":null},{"paper":"/paper/vicon-vision-in-context-operator-networks-for","slug":"vicon-vision-in-context-operator-networks-for","title":"VICON: Vision In-Context Operator Networks for Multi-Physics Fluid Dynamics Prediction","date":"2024-11-25","arxiv_id":"2411.16063","n_code_links":1,"syntology":null},{"paper":null,"slug":"vires-video-instance-repainting-with-sketch","title":"VIRES: Video Instance Repainting via Sketch and Text Guided Generation","date":"2024-11-25","arxiv_id":"2411.16199","n_code_links":0,"syntology":null},{"paper":null,"slug":"vq-sgen-a-vector-quantized-stroke","title":"VQ-SGen: A Vector Quantized Stroke Representation for Creative Sketch Generation","date":"2024-11-25","arxiv_id":"2411.16446","n_code_links":0,"syntology":null},{"paper":null,"slug":"wtdun-wavelet-tree-structured-sampling-and","title":"WTDUN: Wavelet Tree-Structured Sampling and Deep Unfolding Network for Image Compressed Sensing","date":"2024-11-25","arxiv_id":"2411.16336","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-general-sensing-assisted-channel-estimation","title":"A General Sensing-assisted Channel Estimation Framework in Distributed MIMO Network","date":"2024-11-24","arxiv_id":"2411.15995","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-method-for-building-large-language-models","title":"A Method for Building Large Language Models with Predefined KV Cache Capacity","date":"2024-11-24","arxiv_id":"2411.15785","n_code_links":0,"syntology":null},{"paper":"/paper/beyond-adaptive-gradient-fast-controlled","slug":"beyond-adaptive-gradient-fast-controlled","title":"Beyond adaptive gradient: Fast-Controlled Minibatch Algorithm for large-scale optimization","date":"2024-11-24","arxiv_id":"2411.15795","n_code_links":1,"syntology":null},{"paper":"/paper/development-of-pre-trained-transformer-based","slug":"development-of-pre-trained-transformer-based","title":"Development of Pre-Trained Transformer-based Models for the Nepali Language","date":"2024-11-24","arxiv_id":"2411.15734","n_code_links":0,"syntology":null},{"paper":null,"slug":"fasttracktr-towards-fast-multi-object","title":"FastTrackTr:Towards Fast Multi-Object Tracking with Transformers","date":"2024-11-24","arxiv_id":"2411.15811","n_code_links":0,"syntology":null},{"paper":null,"slug":"fixing-the-perspective-a-critical-examination","title":"Fixing the Perspective: A Critical Examination of Zero-1-to-3","date":"2024-11-24","arxiv_id":"2411.15706","n_code_links":0,"syntology":null},{"paper":null,"slug":"gradient-norm-regularization-second-order","title":"Gradient Norm Regularization Second-Order Algorithms for Solving Nonconvex-Strongly Concave Minimax Problems","date":"2024-11-24","arxiv_id":"2411.15769","n_code_links":0,"syntology":null},{"paper":null,"slug":"investigating-factuality-in-long-form-text","title":"Investigating Factuality in Long-Form Text Generation: The Roles of Self-Known and Self-Unknown","date":"2024-11-24","arxiv_id":"2411.15993","n_code_links":0,"syntology":null},{"paper":null,"slug":"letstalk-latent-diffusion-transformer-for","title":"LetsTalk: Latent Diffusion Transformer for Talking Video Synthesis","date":"2024-11-24","arxiv_id":"2411.16748","n_code_links":0,"syntology":null},{"paper":"/paper/llama-moe-v2-exploring-sparsity-of-llama-from","slug":"llama-moe-v2-exploring-sparsity-of-llama-from","title":"LLaMA-MoE v2: Exploring Sparsity of LLaMA from Perspective of Mixture-of-Experts with Post-Training","date":"2024-11-24","arxiv_id":"2411.15708","n_code_links":1,"syntology":{"ran":2,"of":2,"n_ran_checked":0,"n_instrument":2,"unverified":0,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","official":{"repos":["opensparsellms/llama-moe-v2"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"ltcf-net-a-transformer-enhanced-dual-channel","title":"LTCF-Net: A Transformer-Enhanced Dual-Channel Fourier Framework for Low-Light Image Restoration","date":"2024-11-24","arxiv_id":"2411.15740","n_code_links":0,"syntology":null},{"paper":"/paper/medical-slice-transformer-improved-diagnosis","slug":"medical-slice-transformer-improved-diagnosis","title":"Medical Slice Transformer: Improved Diagnosis and Explainability on 3D Medical Images with DINOv2","date":"2024-11-24","arxiv_id":"2411.15802","n_code_links":1,"syntology":null},{"paper":"/paper/nimbus-secure-and-efficient-two-party","slug":"nimbus-secure-and-efficient-two-party","title":"Nimbus: Secure and Efficient Two-Party Inference for Transformers","date":"2024-11-24","arxiv_id":"2411.15707","n_code_links":1,"syntology":{"ran":0,"of":5,"n_ran_checked":0,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"0 ran · 5 unverified","official":{"repos":["secretflow/spu"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":5,"ran_from_kinds":[]}}},{"paper":null,"slug":"pr-mim-delving-deeper-into-partial","title":"PR-MIM: Delving Deeper into Partial Reconstruction in Masked Image Modeling","date":"2024-11-24","arxiv_id":"2411.15746","n_code_links":0,"syntology":null},{"paper":null,"slug":"ramie-retrieval-augmented-multi-task","title":"RAMIE: Retrieval-Augmented Multi-task Information Extraction with Large Language Models on Dietary Supplements","date":"2024-11-24","arxiv_id":"2411.15700","n_code_links":0,"syntology":null},{"paper":"/paper/resclip-residual-attention-for-training-free","slug":"resclip-residual-attention-for-training-free","title":"ResCLIP: Residual Attention for Training-free Dense Vision-language Inference","date":"2024-11-24","arxiv_id":"2411.15851","n_code_links":1,"syntology":null},{"paper":"/paper/self-calibrated-clip-for-training-free-open","slug":"self-calibrated-clip-for-training-free-open","title":"Self-Calibrated CLIP for Training-Free Open-Vocabulary Segmentation","date":"2024-11-24","arxiv_id":"2411.15869","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":2,"n_instrument":3,"unverified":2,"pointer_only":2,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 3 where Syntology's instrument failed) · 2 unverified","official":{"repos":["sulebai/sc-clip"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/test-time-alignment-enhanced-adapter-for","slug":"test-time-alignment-enhanced-adapter-for","title":"Test-time Alignment-Enhanced Adapter for Vision-Language Models","date":"2024-11-24","arxiv_id":"2411.15735","n_code_links":1,"syntology":null},{"paper":null,"slug":"text-guided-coarse-to-fine-fusion-network-for","title":"Text-Guided Coarse-to-Fine Fusion Network for Robust Remote Sensing Visual Question Answering","date":"2024-11-24","arxiv_id":"2411.15770","n_code_links":0,"syntology":null},{"paper":null,"slug":"transfair-transferring-fairness-from-ocular","title":"TransFair: Transferring Fairness from Ocular Disease Classification to Progression Prediction","date":"2024-11-24","arxiv_id":"2412.00051","n_code_links":0,"syntology":null},{"paper":"/paper/a-comparative-analysis-of-transformer-and","slug":"a-comparative-analysis-of-transformer-and","title":"A Comparative Analysis of Transformer and LSTM Models for Detecting Suicidal Ideation on Reddit","date":"2024-11-23","arxiv_id":"2411.15404","n_code_links":1,"syntology":null},{"paper":"/paper/all-that-glitters-approaches-to-evaluations","slug":"all-that-glitters-approaches-to-evaluations","title":"\"All that Glitters\": Approaches to Evaluations with Unreliable Model and Human Annotations","date":"2024-11-23","arxiv_id":"2411.15634","n_code_links":1,"syntology":null},{"paper":null,"slug":"best-of-both-worlds-advantages-of-hybrid","title":"Best of Both Worlds: Advantages of Hybrid Graph Sequence Models","date":"2024-11-23","arxiv_id":"2411.15671","n_code_links":0,"syntology":null},{"paper":null,"slug":"chatbci-a-p300-speller-bci-leveraging-large","title":"ChatBCI: A P300 Speller BCI Leveraging Large Language Models for Improved Sentence Composition in Realistic Scenarios","date":"2024-11-23","arxiv_id":"2411.15395","n_code_links":0,"syntology":null},{"paper":null,"slug":"circuit-design-in-biology-and-machine-1","title":"Circuit design in biology and machine learning. II. Anomaly detection","date":"2024-11-23","arxiv_id":"2411.15647","n_code_links":0,"syntology":null},{"paper":"/paper/devils-in-middle-layers-of-large-vision","slug":"devils-in-middle-layers-of-large-vision","title":"Devils in Middle Layers of Large Vision-Language Models: Interpreting, Detecting and Mitigating Object Hallucinations via Attention Lens","date":"2024-11-23","arxiv_id":"2411.16724","n_code_links":1,"syntology":{"ran":8,"of":10,"n_ran_checked":3,"n_instrument":5,"unverified":2,"pointer_only":10,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 5 where Syntology's instrument failed) · 2 unverified","official":{"repos":["zhangqijiang07/middle_layers_indicating_hallucinations"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"enhancing-instruction-following-capability-of","title":"Enhancing Instruction-Following Capability of Visual-Language Models by Reducing Image Redundancy","date":"2024-11-23","arxiv_id":"2411.15453","n_code_links":0,"syntology":null},{"paper":"/paper/federated-learning-in-chemical-engineering-a","slug":"federated-learning-in-chemical-engineering-a","title":"Federated Learning in Chemical Engineering: A Tutorial on a Framework for Privacy-Preserving Collaboration Across Distributed Data Sources","date":"2024-11-23","arxiv_id":"2411.16737","n_code_links":1,"syntology":null},{"paper":null,"slug":"federated-pca-and-estimation-for-spiked","title":"Federated PCA and Estimation for Spiked Covariance Matrices: Optimal Rates and Efficient Algorithm","date":"2024-11-23","arxiv_id":"2411.15660","n_code_links":0,"syntology":null},{"paper":"/paper/fg-cxr-a-radiologist-aligned-gaze-dataset-for","slug":"fg-cxr-a-radiologist-aligned-gaze-dataset-for","title":"FG-CXR: A Radiologist-Aligned Gaze Dataset for Enhancing Interpretability in Chest X-Ray Report Generation","date":"2024-11-23","arxiv_id":"2411.15413","n_code_links":1,"syntology":{"ran":1,"of":5,"n_ran_checked":0,"n_instrument":1,"unverified":4,"pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 4 unverified","official":{"repos":["uark-aicv/fg-cxr"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"freepruner-a-training-free-approach-for-large","title":"freePruner: A Training-free Approach for Large Multimodal Model Acceleration","date":"2024-11-23","arxiv_id":"2411.15446","n_code_links":0,"syntology":null},{"paper":"/paper/geoai-enhanced-community-detection-on-spatial","slug":"geoai-enhanced-community-detection-on-spatial","title":"GeoAI-Enhanced Community Detection on Spatial Networks with Graph Deep Learning","date":"2024-11-23","arxiv_id":"2411.15428","n_code_links":1,"syntology":null},{"paper":null,"slug":"improving-next-tokens-via-second-last","title":"Improving Next Tokens via Second-Last Predictions with Generate and Refine","date":"2024-11-23","arxiv_id":"2411.15661","n_code_links":0,"syntology":null},{"paper":null,"slug":"inducing-human-like-biases-in-moral-reasoning","title":"Inducing Human-like Biases in Moral Reasoning Language Models","date":"2024-11-23","arxiv_id":"2411.15386","n_code_links":0,"syntology":null},{"paper":"/paper/large-scale-text-to-image-model-with","slug":"large-scale-text-to-image-model-with","title":"Large-Scale Text-to-Image Model with Inpainting is a Zero-Shot Subject-Driven Image Generator","date":"2024-11-23","arxiv_id":"2411.15466","n_code_links":1,"syntology":null},{"paper":"/paper/semantic-shield-defending-vision-language-1","slug":"semantic-shield-defending-vision-language-1","title":"Semantic Shield: Defending Vision-Language Models Against Backdooring and Poisoning via Fine-grained Knowledge Alignment","date":"2024-11-23","arxiv_id":"2411.15673","n_code_links":1,"syntology":null},{"paper":"/paper/sprint-enables-interpretable-and-ultra-fast","slug":"sprint-enables-interpretable-and-ultra-fast","title":"Scaling Structure Aware Virtual Screening to Billions of Molecules with SPRINT","date":"2024-11-23","arxiv_id":"2411.15418","n_code_links":1,"syntology":null},{"paper":"/paper/tangnn-a-concise-scalable-and-effective-graph","slug":"tangnn-a-concise-scalable-and-effective-graph","title":"TANGNN: a Concise, Scalable and Effective Graph Neural Networks with Top-m Attention Mechanism for Graph Representation Learning","date":"2024-11-23","arxiv_id":"2411.15458","n_code_links":1,"syntology":null},{"paper":"/paper/towards-satellite-image-road-graph-extraction","slug":"towards-satellite-image-road-graph-extraction","title":"Towards Satellite Image Road Graph Extraction: A Global-Scale Dataset and A Novel Method","date":"2024-11-23","arxiv_id":"2411.16733","n_code_links":1,"syntology":null},{"paper":null,"slug":"traditional-chinese-medicine-case-analysis","title":"Traditional Chinese Medicine Case Analysis System for High-Level Semantic Abstraction: Optimized with Prompt and RAG","date":"2024-11-23","arxiv_id":"2411.15491","n_code_links":0,"syntology":null},{"paper":null,"slug":"training-an-open-vocabulary-monocular-3d","title":"Training an Open-Vocabulary Monocular 3D Object Detection Model without 3D Data","date":"2024-11-23","arxiv_id":"2411.15657","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-real-time-detr-approach-to-bangladesh-road","title":"A Real-Time DETR Approach to Bangladesh Road Object Detection for Autonomous Vehicles","date":"2024-11-22","arxiv_id":"2411.15110","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-attention-based-framework-for-fair","title":"An Attention-based Framework for Fair Contrastive Learning","date":"2024-11-22","arxiv_id":"2411.14765","n_code_links":0,"syntology":null},{"paper":null,"slug":"astro-hep-bert-a-bidirectional-language-model","title":"Astro-HEP-BERT: A bidirectional language model for studying the meanings of concepts in astrophysics and high energy physics","date":"2024-11-22","arxiv_id":"2411.14877","n_code_links":0,"syntology":null},{"paper":null,"slug":"boundless-across-domains-a-new-paradigm-of","title":"Boundless Across Domains: A New Paradigm of Adaptive Feature and Cross-Attention for Domain Generalization in Medical Image Segmentation","date":"2024-11-22","arxiv_id":"2411.14883","n_code_links":0,"syntology":null},{"paper":null,"slug":"comparative-analysis-of-pooling-mechanisms-in","title":"Comparative Analysis of Pooling Mechanisms in LLMs: A Sentiment Analysis Perspective","date":"2024-11-22","arxiv_id":"2411.14654","n_code_links":0,"syntology":null},{"paper":"/paper/cross-group-attention-and-group-wise-rolling","slug":"cross-group-attention-and-group-wise-rolling","title":"Learning Modality-Aware Representations: Adaptive Group-wise Interaction Network for Multimodal MRI Synthesis","date":"2024-11-22","arxiv_id":"2411.14684","n_code_links":1,"syntology":null},{"paper":null,"slug":"cross-modal-pre-aligned-method-with-global","title":"Cross-Modal Pre-Aligned Method with Global and Local Information for Remote-Sensing Image and Text Retrieval","date":"2024-11-22","arxiv_id":"2411.14704","n_code_links":0,"syntology":null},{"paper":null,"slug":"defective-edge-detection-using-cascaded","title":"Defective Edge Detection Using Cascaded Ensemble Canny Operator","date":"2024-11-22","arxiv_id":"2411.14868","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-visual-triggers-in-cannabis-imagery","title":"Detecting Visual Triggers in Cannabis Imagery: A CLIP-Based Multi-Labeling Framework with Local-Global Aggregation","date":"2024-11-22","arxiv_id":"2412.08648","n_code_links":0,"syntology":null},{"paper":null,"slug":"don-t-mesh-with-me-generating-constructive","title":"Don't Mesh with Me: Generating Constructive Solid Geometry Instead of Meshes by Fine-Tuning a Code-Generation LLM","date":"2024-11-22","arxiv_id":"2411.15279","n_code_links":0,"syntology":null},{"paper":"/paper/efficientvim-efficient-vision-mamba-with","slug":"efficientvim-efficient-vision-mamba-with","title":"EfficientViM: Efficient Vision Mamba with Hidden State Mixer based State Space Duality","date":"2024-11-22","arxiv_id":"2411.15241","n_code_links":2,"syntology":null},{"paper":null,"slug":"elastiformer-learned-redundancy-reduction-in","title":"ElastiFormer: Learned Redundancy Reduction in Transformer via Self-Distillation","date":"2024-11-22","arxiv_id":"2411.15281","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-vision-transformer-models-for","slug":"evaluating-vision-transformer-models-for","title":"Evaluating Vision Transformer Models for Visual Quality Control in Industrial Manufacturing","date":"2024-11-22","arxiv_id":"2411.14953","n_code_links":1,"syntology":null},{"paper":null,"slug":"exploring-the-use-of-machine-learning-weather","title":"Exploring the Use of Machine Learning Weather Models in Data Assimilation","date":"2024-11-22","arxiv_id":"2411.14677","n_code_links":0,"syntology":null},{"paper":null,"slug":"fast-high-quality-enhanced-imaging-algorithm","title":"Fast High-Quality Enhanced Imaging Algorithm for Layered Dielectric Targets Based on MMW MIMO-SAR System","date":"2024-11-22","arxiv_id":"2411.14837","n_code_links":0,"syntology":null},{"paper":null,"slug":"headrouter-a-training-free-image-editing","title":"HeadRouter: A Training-free Image Editing Framework for MM-DiTs by Adaptively Routing Attention Heads","date":"2024-11-22","arxiv_id":"2411.15034","n_code_links":0,"syntology":null},{"paper":null,"slug":"ict-image-object-cross-level-trusted","title":"ICT: Image-Object Cross-Level Trusted Intervention for Mitigating Object Hallucination in Large Vision-Language Models","date":"2024-11-22","arxiv_id":"2411.15268","n_code_links":0,"syntology":null},{"paper":"/paper/is-attention-all-you-need-for-actigraphy","slug":"is-attention-all-you-need-for-actigraphy","title":"AI Foundation Models for Wearable Movement Data in Mental Health Research","date":"2024-11-22","arxiv_id":"2411.15240","n_code_links":1,"syntology":null},{"paper":null,"slug":"j-invariant-volume-shuffle-for-self","title":"J-Invariant Volume Shuffle for Self-Supervised Cryo-Electron Tomogram Denoising on Single Noisy Volume","date":"2024-11-22","arxiv_id":"2411.15248","n_code_links":0,"syntology":null},{"paper":"/paper/kbada-efficient-self-adaptation-on-specific","slug":"kbada-efficient-self-adaptation-on-specific","title":"KBAlign: Efficient Self Adaptation on Specific Knowledge Bases","date":"2024-11-22","arxiv_id":"2411.14790","n_code_links":1,"syntology":null},{"paper":"/paper/mme-survey-a-comprehensive-survey-on","slug":"mme-survey-a-comprehensive-survey-on","title":"MME-Survey: A Comprehensive Survey on Evaluation of Multimodal LLMs","date":"2024-11-22","arxiv_id":"2411.15296","n_code_links":4,"syntology":null},{"paper":"/paper/multi-granularity-interest-retrieval-and","slug":"multi-granularity-interest-retrieval-and","title":"Multi-granularity Interest Retrieval and Refinement Network for Long-Term User Behavior Modeling in CTR Prediction","date":"2024-11-22","arxiv_id":"2411.15005","n_code_links":3,"syntology":null},{"paper":"/paper/multiset-transformer-advancing-representation","slug":"multiset-transformer-advancing-representation","title":"Multiset Transformer: Advancing Representation Learning in Persistence Diagrams","date":"2024-11-22","arxiv_id":"2411.14662","n_code_links":1,"syntology":null},{"paper":"/paper/ominicontrol-minimal-and-universal-control","slug":"ominicontrol-minimal-and-universal-control","title":"OminiControl: Minimal and Universal Control for Diffusion Transformer","date":"2024-11-22","arxiv_id":"2411.15098","n_code_links":2,"syntology":{"ran":1,"of":2,"n_ran_checked":1,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["Yuanshi9815/OminiControl","Yuanshi9815/Subjects200K"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"point-cloud-understanding-via-attention","title":"Point Cloud Understanding via Attention-Driven Contrastive Learning","date":"2024-11-22","arxiv_id":"2411.14744","n_code_links":0,"syntology":null},{"paper":null,"slug":"purrfessor-a-fine-tuned-multimodal-llava-diet","title":"Purrfessor: A Fine-tuned Multimodal LLaVA Diet Health Chatbot","date":"2024-11-22","arxiv_id":"2411.14925","n_code_links":0,"syntology":null},{"paper":"/paper/recursive-gaussian-process-state-space-model","slug":"recursive-gaussian-process-state-space-model","title":"Recursive Gaussian Process State Space Model","date":"2024-11-22","arxiv_id":"2411.14679","n_code_links":2,"syntology":null},{"paper":null,"slug":"red-effective-trajectory-representation","title":"RED: Effective Trajectory Representation Learning with Comprehensive Information","date":"2024-11-22","arxiv_id":"2411.15096","n_code_links":0,"syntology":null},{"paper":null,"slug":"resolution-agnostic-transformer-based-climate","title":"Resolution-Agnostic Transformer-based Climate Downscaling","date":"2024-11-22","arxiv_id":"2411.14774","n_code_links":0,"syntology":null},{"paper":null,"slug":"safelight-enhancing-security-in-optical","title":"SafeLight: Enhancing Security in Optical Convolutional Neural Network Accelerators","date":"2024-11-22","arxiv_id":"2411.16712","n_code_links":0,"syntology":null},{"paper":"/paper/scribeagent-towards-specialized-web-agents","slug":"scribeagent-towards-specialized-web-agents","title":"ScribeAgent: Towards Specialized Web Agents Using Production-Scale Workflow Data","date":"2024-11-22","arxiv_id":"2411.15004","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["colonylabs/ScribeAgent"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"simplifying-clip-unleashing-the-power-of","title":"Simplifying CLIP: Unleashing the Power of Large-Scale Models on Consumer-level Computers","date":"2024-11-22","arxiv_id":"2411.14789","n_code_links":0,"syntology":null},{"paper":"/paper/texgen-a-generative-diffusion-model-for-mesh","slug":"texgen-a-generative-diffusion-model-for-mesh","title":"TEXGen: a Generative Diffusion Model for Mesh Textures","date":"2024-11-22","arxiv_id":"2411.14740","n_code_links":1,"syntology":null},{"paper":null,"slug":"transforming-nlu-with-babylon-a-case-study-in","title":"Transforming NLU with Babylon: A Case Study in Development of Real-time, Edge-Efficient, Multi-Intent Translation System for Automated Drive-Thru Ordering","date":"2024-11-22","arxiv_id":"2411.15372","n_code_links":0,"syntology":null},{"paper":null,"slug":"when-spatial-meets-temporal-in-action","title":"When Spatial meets Temporal in Action Recognition","date":"2024-11-22","arxiv_id":"2411.15284","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-accuracy-improving-method-for-advertising","title":"An accuracy improving method for advertising click through rate prediction based on enhanced xDeepFM model","date":"2024-11-21","arxiv_id":"2411.15223","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-experimental-study-on-data-augmentation","title":"An Experimental Study on Data Augmentation Techniques for Named Entity Recognition on Low-Resource Domains","date":"2024-11-21","arxiv_id":"2411.14551","n_code_links":0,"syntology":null},{"paper":null,"slug":"assessment-of-llm-responses-to-end-user","title":"Assessment of LLM Responses to End-user Security Questions","date":"2024-11-21","arxiv_id":"2411.14571","n_code_links":0,"syntology":null},{"paper":"/paper/benchmarking-gpt-4-against-human-translators","slug":"benchmarking-gpt-4-against-human-translators","title":"Benchmarking GPT-4 against Human Translators: A Comprehensive Evaluation Across Languages, Domains, and Expertise Levels","date":"2024-11-21","arxiv_id":"2411.13775","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["elliottyan/gpt_versus_mt_experts"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":"/paper/bert-based-approach-for-automating-course","slug":"bert-based-approach-for-automating-course","title":"BERT-Based Approach for Automating Course Articulation Matrix Construction with Explainable AI","date":"2024-11-21","arxiv_id":"2411.14254","n_code_links":1,"syntology":null},{"paper":"/paper/cliper-hierarchically-improving-spatial","slug":"cliper-hierarchically-improving-spatial","title":"CLIPer: Hierarchically Improving Spatial Representation of CLIP for Open-Vocabulary Semantic Segmentation","date":"2024-11-21","arxiv_id":"2411.13836","n_code_links":1,"syntology":null},{"paper":null,"slug":"contrasting-local-and-global-modeling-with","title":"Contrasting local and global modeling with machine learning and satellite data: A case study estimating tree canopy height in African savannas","date":"2024-11-21","arxiv_id":"2411.14354","n_code_links":0,"syntology":null},{"paper":"/paper/dino-x-a-unified-vision-model-for-open-world","slug":"dino-x-a-unified-vision-model-for-open-world","title":"DINO-X: A Unified Vision Model for Open-World Object Detection and Understanding","date":"2024-11-21","arxiv_id":"2411.14347","n_code_links":1,"syntology":null},{"paper":null,"slug":"do-i-know-this-entity-knowledge-awareness-and","title":"Do I Know This Entity? Knowledge Awareness and Hallucinations in Language Models","date":"2024-11-21","arxiv_id":"2411.14257","n_code_links":0,"syntology":null},{"paper":"/paper/evaluating-the-robustness-of-analogical","slug":"evaluating-the-robustness-of-analogical","title":"Evaluating the Robustness of Analogical Reasoning in Large Language Models","date":"2024-11-21","arxiv_id":"2411.14215","n_code_links":1,"syntology":{"ran":1,"of":1,"n_ran_checked":1,"n_instrument":0,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["marthaflinderslewis/robust-analogy"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"explaining-gpt-4-s-schema-of-depression-using","title":"Explaining GPT-4's Schema of Depression Using Machine Behavior Analysis","date":"2024-11-21","arxiv_id":"2411.13800","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-applications-of-topological-data","slug":"exploring-applications-of-topological-data","title":"Exploring applications of topological data analysis in stock index movement prediction","date":"2024-11-21","arxiv_id":"2411.13881","n_code_links":1,"syntology":null},{"paper":null,"slug":"fastrag-retrieval-augmented-generation-for","title":"FastRAG: Retrieval Augmented Generation for Semi-structured Data","date":"2024-11-21","arxiv_id":"2411.13773","n_code_links":0,"syntology":null},{"paper":"/paper/g-rag-knowledge-expansion-in-material-science","slug":"g-rag-knowledge-expansion-in-material-science","title":"G-RAG: Knowledge Expansion in Material Science","date":"2024-11-21","arxiv_id":"2411.14592","n_code_links":1,"syntology":null},{"paper":null,"slug":"generative-fuzzy-system-for-sequence","title":"Generative Fuzzy System for Sequence Generation","date":"2024-11-21","arxiv_id":"2411.13867","n_code_links":0,"syntology":null},{"paper":"/paper/global-and-local-attention-based-transformer","slug":"global-and-local-attention-based-transformer","title":"Global and Local Attention-Based Transformer for Hyperspectral Image Change Detection","date":"2024-11-21","arxiv_id":"2411.14109","n_code_links":1,"syntology":null},{"paper":"/paper/gmai-vl-gmai-vl-5-5m-a-large-vision-language","slug":"gmai-vl-gmai-vl-5-5m-a-large-vision-language","title":"GMAI-VL & GMAI-VL-5.5M: A Large Vision-Language Model and A Comprehensive Multimodal Dataset Towards General Medical AI","date":"2024-11-21","arxiv_id":"2411.14522","n_code_links":1,"syntology":{"ran":10,"of":10,"n_ran_checked":9,"n_instrument":1,"unverified":0,"pointer_only":2,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["uni-medical/gmai-vl"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"graph-domain-adaptation-with-dual-branch","title":"Graph Domain Adaptation with Dual-branch Encoder and Two-level Alignment for Whole Slide Image-based Survival Prediction","date":"2024-11-21","arxiv_id":"2411.14001","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-from-silly-questions-improves-large","title":"Learning from \"Silly\" Questions Improves Large Language Models, But Only Slightly","date":"2024-11-21","arxiv_id":"2411.14121","n_code_links":0,"syntology":null}],"record_sha256":"85f5655b8f322acc2e6d27a9c6824c647ff7c69dd40239f46e1ca62b6d6032ff","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}