{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/attention/papers/216","list_of":"/method/attention","method":"Attention","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":216,"pages_in_order":316,"rows_per_page":100,"rows":[21501,21600],"of":31583,"counts":{"archive_papers_tagged":31583,"with_a_code_link":13473,"where_syntology_ran_a_sample":3998,"not_listed_spam_title":0,"listed":31583,"listed_where_code_ran":3998,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3366,"every_run_a_failure_of_syntologys_instrument":632,"listed_with_a_run_with_no_instrument_failure":3366,"listed_every_run_a_failure_of_syntologys_instrument":632,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/attention","prev":"/method/attention/papers/215","next":"/method/attention/papers/217","papers":[{"paper":"/paper/sat-size-aware-transformer-for-3d-point-cloud","slug":"sat-size-aware-transformer-for-3d-point-cloud","title":"SAT: Size-Aware Transformer for 3D Point Cloud Semantic Segmentation","date":"2023-01-17","arxiv_id":"2301.06869","n_code_links":0,"syntology":null},{"paper":"/paper/swindepth-unsupervised-depth-estimation-using","slug":"swindepth-unsupervised-depth-estimation-using","title":"SwinDepth: Unsupervised Depth Estimation using Monocular Sequences via Swin Transformer and Densely Cascaded Network","date":"2023-01-17","arxiv_id":"2301.06715","n_code_links":1,"syntology":null},{"paper":"/paper/tracing-and-manipulating-intermediate-values","slug":"tracing-and-manipulating-intermediate-values","title":"Tracing and Manipulating Intermediate Values in Neural Math Problem Solvers","date":"2023-01-17","arxiv_id":"2301.06758","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformer-based-implementation-for","title":"Transformer Based Implementation for Automatic Book Summarization","date":"2023-01-17","arxiv_id":"2301.07057","n_code_links":0,"syntology":null},{"paper":"/paper/an-error-guided-correction-model-for-chinese","slug":"an-error-guided-correction-model-for-chinese","title":"An Error-Guided Correction Model for Chinese Spelling Error Correction","date":"2023-01-16","arxiv_id":"2301.06323","n_code_links":1,"syntology":null},{"paper":null,"slug":"bayesspeech-a-bayesian-transformer-network","title":"BayesSpeech: A Bayesian Transformer Network for Automatic Speech Recognition","date":"2023-01-16","arxiv_id":"2301.11276","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-target-speaker-extraction-with","title":"Improving Target Speaker Extraction with Sparse LDA-transformed Speaker Embeddings","date":"2023-01-16","arxiv_id":"2301.06277","n_code_links":0,"syntology":null},{"paper":"/paper/tdstf-transformer-based-diffusion","slug":"tdstf-transformer-based-diffusion","title":"A Transformer-based Diffusion Probabilistic Model for Heart Rate and Blood Pressure Forecasting in Intensive Care Unit","date":"2023-01-16","arxiv_id":"2301.06625","n_code_links":1,"syntology":null},{"paper":"/paper/tedb-system-description-to-a-shared-task-on","slug":"tedb-system-description-to-a-shared-task-on","title":"TEDB System Description to a Shared Task on Euphemism Detection 2022","date":"2023-01-16","arxiv_id":"2301.06602","n_code_links":1,"syntology":null},{"paper":"/paper/dsvt-dynamic-sparse-voxel-transformer-with","slug":"dsvt-dynamic-sparse-voxel-transformer-with","title":"DSVT: Dynamic Sparse Voxel Transformer with Rotated Sets","date":"2023-01-15","arxiv_id":"2301.06051","n_code_links":4,"syntology":null},{"paper":null,"slug":"improving-noise-robustness-for-spoken-content","title":"Improving Noise Robustness for Spoken Content Retrieval using Semi-supervised ASR and N-best Transcripts for BERT-based Ranking Models","date":"2023-01-15","arxiv_id":"2301.06056","n_code_links":0,"syntology":null},{"paper":"/paper/t2m-gpt-generating-human-motion-from-textual","slug":"t2m-gpt-generating-human-motion-from-textual","title":"T2M-GPT: Generating Human Motion from Textual Descriptions with Discrete Representations","date":"2023-01-15","arxiv_id":"2301.06052","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 4 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["Mael-zys/T2M-GPT"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":4,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"salient-sign-detection-in-safe-autonomous","title":"Salient Sign Detection In Safe Autonomous Driving: AI Which Reasons Over Full Visual Context","date":"2023-01-14","arxiv_id":"2301.05804","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-speech-and-text-based","title":"Automated speech- and text-based classification of neuropsychiatric conditions in a multidiagnostic setting","date":"2023-01-13","arxiv_id":"2301.06916","n_code_links":0,"syntology":null},{"paper":"/paper/efficient-activation-function-optimization","slug":"efficient-activation-function-optimization","title":"Efficient Activation Function Optimization through Surrogate Modeling","date":"2023-01-13","arxiv_id":"2301.05785","n_code_links":2,"syntology":{"ran":3,"of":3,"n_ran_checked":3,"n_instrument":0,"unverified":0,"pointer_only":3,"phrase":"3 ran (of which 1 constructed an object rather than computing a result; 3 with no instrument failure: 2 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["cognizant-ai-labs/aquasurf","cognizant-ai-labs/act-bench"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":1,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"text-to-point-cloud-localization-with","title":"Text to Point Cloud Localization with Relation-Enhanced Transformer","date":"2023-01-13","arxiv_id":"2301.05372","n_code_links":0,"syntology":null},{"paper":"/paper/adversarial-adaptation-for-french-named","slug":"adversarial-adaptation-for-french-named","title":"Adversarial Adaptation for French Named Entity Recognition","date":"2023-01-12","arxiv_id":"2301.05220","n_code_links":1,"syntology":null},{"paper":"/paper/vits-for-sits-vision-transformers-for","slug":"vits-for-sits-vision-transformers-for","title":"ViTs for SITS: Vision Transformers for Satellite Image Time Series","date":"2023-01-12","arxiv_id":"2301.04944","n_code_links":3,"syntology":{"ran":11,"of":11,"n_ran_checked":11,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"11 ran (of which 3 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["michaeltrs/deepsatmodels"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/adapointr-diverse-point-cloud-completion-with","slug":"adapointr-diverse-point-cloud-completion-with","title":"AdaPoinTr: Diverse Point Cloud Completion with Adaptive Geometry-Aware Transformers","date":"2023-01-11","arxiv_id":"2301.04545","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":2,"n_instrument":2,"unverified":2,"pointer_only":3,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["yuxumin/PoinTr"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"anomalies-representations-and-self","title":"Anomalies, Representations, and Self-Supervision","date":"2023-01-11","arxiv_id":"2301.04660","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-model-with-attention-mechanism","title":"Super-resolution of Ray-tracing Channel Simulation via Attention Mechanism based Deep Learning Model","date":"2023-01-11","arxiv_id":"2301.04479","n_code_links":0,"syntology":null},{"paper":null,"slug":"dynamic-background-reconstruction-via","title":"Dynamic Background Reconstruction via MAE for Infrared Small Target Detection","date":"2023-01-11","arxiv_id":"2301.04497","n_code_links":0,"syntology":null},{"paper":"/paper/gpt-as-knowledge-worker-a-zero-shot","slug":"gpt-as-knowledge-worker-a-zero-shot","title":"GPT as Knowledge Worker: A Zero-Shot Evaluation of (AI)CPA Capabilities","date":"2023-01-11","arxiv_id":"2301.04408","n_code_links":1,"syntology":null},{"paper":"/paper/head-free-lightweight-semantic-segmentation","slug":"head-free-lightweight-semantic-segmentation","title":"Head-Free Lightweight Semantic Segmentation with Linear Transformer","date":"2023-01-11","arxiv_id":"2301.04648","n_code_links":1,"syntology":null},{"paper":"/paper/narrowbert-accelerating-masked-language-model","slug":"narrowbert-accelerating-masked-language-model","title":"NarrowBERT: Accelerating Masked Language Model Pretraining and Inference","date":"2023-01-11","arxiv_id":"2301.04761","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"0 ran · 2 unverified","official":{"repos":["lihaoxin2020/narrowbert"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":null,"slug":"topics-in-contextualised-attention-embeddings","title":"Topics in Contextualised Attention Embeddings","date":"2023-01-11","arxiv_id":"2301.04339","n_code_links":0,"syntology":null},{"paper":"/paper/dynamic-grained-encoder-for-vision-1","slug":"dynamic-grained-encoder-for-vision-1","title":"Dynamic Grained Encoder for Vision Transformers","date":"2023-01-10","arxiv_id":"2301.03831","n_code_links":1,"syntology":{"ran":1,"of":5,"n_ran_checked":1,"n_instrument":0,"unverified":4,"pointer_only":5,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 4 unverified","official":{"repos":["stevengrove/vtpack"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"language-models-sounds-the-death-knell-of","title":"Language Models sounds the Death Knell of Knowledge Graphs","date":"2023-01-10","arxiv_id":"2301.03980","n_code_links":0,"syntology":null},{"paper":"/paper/predicting-hateful-discussions-on-reddit","slug":"predicting-hateful-discussions-on-reddit","title":"Predicting Hateful Discussions on Reddit using Graph Transformer Networks and Communal Context","date":"2023-01-10","arxiv_id":"2301.04248","n_code_links":1,"syntology":null},{"paper":null,"slug":"recommending-root-cause-and-mitigation-steps","title":"Recommending Root-Cause and Mitigation Steps for Cloud Incidents using Large Language Models","date":"2023-01-10","arxiv_id":"2301.03797","n_code_links":0,"syntology":null},{"paper":null,"slug":"streaming-punctuation-a-novel-punctuation","title":"Streaming Punctuation: A Novel Punctuation Technique Leveraging Bidirectional Context for Continuous Speech Recognition","date":"2023-01-10","arxiv_id":"2301.03819","n_code_links":0,"syntology":null},{"paper":"/paper/there-is-no-big-brother-or-small-brother","slug":"there-is-no-big-brother-or-small-brother","title":"There is No Big Brother or Small Brother: Knowledge Infusion in Language Models for Link Prediction and Question Answering","date":"2023-01-10","arxiv_id":"2301.04013","n_code_links":2,"syntology":null},{"paper":"/paper/unsupervised-mandarin-cantonese-machine","slug":"unsupervised-mandarin-cantonese-machine","title":"Unsupervised Mandarin-Cantonese Machine Translation","date":"2023-01-10","arxiv_id":"2301.03971","n_code_links":1,"syntology":null},{"paper":"/paper/a-study-on-the-generality-of-neural-network","slug":"a-study-on-the-generality-of-neural-network","title":"A Study on the Generality of Neural Network Structures for Monocular Depth Estimation","date":"2023-01-09","arxiv_id":"2301.03169","n_code_links":1,"syntology":null},{"paper":"/paper/advances-in-medical-image-analysis-with","slug":"advances-in-medical-image-analysis-with","title":"Advances in Medical Image Analysis with Vision Transformers: A Comprehensive Review","date":"2023-01-09","arxiv_id":"2301.03505","n_code_links":1,"syntology":null},{"paper":"/paper/an-impartial-transformer-for-story","slug":"an-impartial-transformer-for-story","title":"An Impartial Transformer for Story Visualization","date":"2023-01-09","arxiv_id":"2301.03563","n_code_links":0,"syntology":null},{"paper":"/paper/demt-deformable-mixer-transformer-for-multi","slug":"demt-deformable-mixer-transformer-for-multi","title":"DeMT: Deformable Mixer Transformer for Multi-Task Learning of Dense Prediction","date":"2023-01-09","arxiv_id":"2301.03461","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":6,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["yangyangxu0/demt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/designing-bert-for-convolutional-networks","slug":"designing-bert-for-convolutional-networks","title":"Designing BERT for Convolutional Networks: Sparse and Hierarchical Masked Modeling","date":"2023-01-09","arxiv_id":"2301.03580","n_code_links":2,"syntology":{"ran":13,"of":14,"n_ran_checked":12,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["keyu-tian/spark"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"logically-at-factify-2023-a-multi-modal-fact","title":"Logically at Factify 2: A Multi-Modal Fact Checking System Based on Evidence Retrieval techniques and Transformer Encoder Architecture","date":"2023-01-09","arxiv_id":"2301.03127","n_code_links":0,"syntology":null},{"paper":null,"slug":"online-fake-review-detection-using-supervised","title":"Online Fake Review Detection Using Supervised Machine Learning And BERT Model","date":"2023-01-09","arxiv_id":"2301.03225","n_code_links":0,"syntology":null},{"paper":null,"slug":"universal-multimodal-representation-for","title":"Universal Multimodal Representation for Language Understanding","date":"2023-01-09","arxiv_id":"2301.03344","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-generation-of-german-drama-texts","title":"Automatic Generation of German Drama Texts Using Fine Tuned GPT-2 Models","date":"2023-01-08","arxiv_id":"2301.03119","n_code_links":0,"syntology":null},{"paper":"/paper/deepmatcher-a-deep-transformer-based-network","slug":"deepmatcher-a-deep-transformer-based-network","title":"DeepMatcher: A Deep Transformer-based Network for Robust and Accurate Local Feature Matching","date":"2023-01-08","arxiv_id":"2301.02993","n_code_links":1,"syntology":null},{"paper":"/paper/hrtransnet-hrformer-driven-two-modality","slug":"hrtransnet-hrformer-driven-two-modality","title":"HRTransNet: HRFormer-Driven Two-Modality Salient Object Detection","date":"2023-01-08","arxiv_id":"2301.03036","n_code_links":1,"syntology":null},{"paper":"/paper/app-review-driven-collaborative-bug-finding","slug":"app-review-driven-collaborative-bug-finding","title":"App Review Driven Collaborative Bug Finding","date":"2023-01-07","arxiv_id":"2301.02818","n_code_links":1,"syntology":null},{"paper":"/paper/rlas-biabc-a-reinforcement-learning-based","slug":"rlas-biabc-a-reinforcement-learning-based","title":"RLAS-BIABC: A Reinforcement Learning-Based Answer Selection Using the BERT Model Boosted by an Improved ABC Algorithm","date":"2023-01-07","arxiv_id":"2301.02807","n_code_links":0,"syntology":null},{"paper":"/paper/codetalker-speech-driven-3d-facial-animation","slug":"codetalker-speech-driven-3d-facial-animation","title":"CodeTalker: Speech-Driven 3D Facial Animation with Discrete Motion Prior","date":"2023-01-06","arxiv_id":"2301.02379","n_code_links":1,"syntology":{"ran":8,"of":13,"n_ran_checked":8,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 1 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["Doubiiu/CodeTalker"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"conditional-generation-of-paired-antibody","title":"Generative Antibody Design for Complementary Chain Pairing Sequences through Encoder-Decoder Language Model","date":"2023-01-06","arxiv_id":"2301.02748","n_code_links":0,"syntology":null},{"paper":null,"slug":"does-compressing-activations-help-model","title":"Does compressing activations help model parallel training?","date":"2023-01-06","arxiv_id":"2301.02654","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-efficient-few-shot-adaptation-for","slug":"exploring-efficient-few-shot-adaptation-for","title":"Exploring Efficient Few-shot Adaptation for Vision Transformers","date":"2023-01-06","arxiv_id":"2301.02419","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-genre-music-transformer-composing-full","title":"Multi-Genre Music Transformer -- Composing Full Length Musical Piece","date":"2023-01-06","arxiv_id":"2301.02385","n_code_links":0,"syntology":null},{"paper":null,"slug":"systems-for-parallel-and-distributed-large","title":"Systems for Parallel and Distributed Large-Model Deep Learning Training","date":"2023-01-06","arxiv_id":"2301.02691","n_code_links":0,"syntology":null},{"paper":null,"slug":"task-aware-feature-extraction-framework-for","title":"Adaptive Pattern Extraction Multi-Task Learning for Multi-Step Conversion Estimations","date":"2023-01-06","arxiv_id":"2301.02494","n_code_links":0,"syntology":null},{"paper":"/paper/adaptively-clustering-neighbor-elements-for","slug":"adaptively-clustering-neighbor-elements-for","title":"Adaptively Clustering Neighbor Elements for Image-Text Generation","date":"2023-01-05","arxiv_id":"2301.01955","n_code_links":1,"syntology":null},{"paper":null,"slug":"cat-localization-and-identification-cascade-1","title":"CAT: LoCalization and IdentificAtion Cascade Detection Transformer for Open-World Object Detection","date":"2023-01-05","arxiv_id":"2301.01970","n_code_links":0,"syntology":null},{"paper":"/paper/critical-perspectives-a-benchmark-revealing","slug":"critical-perspectives-a-benchmark-revealing","title":"Critical Perspectives: A Benchmark Revealing Pitfalls in PerspectiveAPI","date":"2023-01-05","arxiv_id":"2301.01874","n_code_links":1,"syntology":null},{"paper":null,"slug":"enabling-augmented-segmentation-and","title":"Enabling Augmented Segmentation and Registration in Ultrasound-Guided Spinal Surgery via Realistic Ultrasound Synthesis from Diagnostic CT Volume","date":"2023-01-05","arxiv_id":"2301.01940","n_code_links":0,"syntology":null},{"paper":"/paper/learning-feature-recovery-transformer-for","slug":"learning-feature-recovery-transformer-for","title":"Learning Feature Recovery Transformer for Occluded Person Re-identification","date":"2023-01-05","arxiv_id":"2301.01879","n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-trajectory-word-alignments-for-video","title":"Learning Trajectory-Word Alignments for Video-Language Tasks","date":"2023-01-05","arxiv_id":"2301.01953","n_code_links":0,"syntology":null},{"paper":null,"slug":"ms-dino-efficient-distributed-training-of","title":"Single-round Self-supervised Distributed Learning using Vision Transformer","date":"2023-01-05","arxiv_id":"2301.02064","n_code_links":0,"syntology":null},{"paper":null,"slug":"scalable-communication-for-multi-agent","title":"Scalable Communication for Multi-Agent Reinforcement Learning via Transformer-Based Email Mechanism","date":"2023-01-05","arxiv_id":"2301.01919","n_code_links":0,"syntology":null},{"paper":null,"slug":"sequentially-controlled-text-generation-1","title":"Sequentially Controlled Text Generation","date":"2023-01-05","arxiv_id":"2301.02299","n_code_links":0,"syntology":null},{"paper":"/paper/towards-autoformalization-of-mathematics-and","slug":"towards-autoformalization-of-mathematics-and","title":"Towards Autoformalization of Mathematics and Code Correctness: Experiments with Elementary Proofs","date":"2023-01-05","arxiv_id":"2301.02195","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 5 with no instrument failure: 0 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["gc974517/autoformalization"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":"/paper/towards-long-term-time-series-forecasting","slug":"towards-long-term-time-series-forecasting","title":"Towards Long-Term Time-Series Forecasting: Feature, Pattern, and Distribution","date":"2023-01-05","arxiv_id":"2301.02068","n_code_links":1,"syntology":null},{"paper":"/paper/extending-source-code-pre-trained-language","slug":"extending-source-code-pre-trained-language","title":"Extending Source Code Pre-Trained Language Models to Summarise Decompiled Binaries","date":"2023-01-04","arxiv_id":"2301.01701","n_code_links":1,"syntology":null},{"paper":null,"slug":"infomaxformer-maximum-entropy-transformer-for","title":"Infomaxformer: Maximum Entropy Transformer for Long Time-Series Forecasting Problem","date":"2023-01-04","arxiv_id":"2301.01772","n_code_links":0,"syntology":null},{"paper":"/paper/inpars-v2-large-language-models-as-efficient","slug":"inpars-v2-large-language-models-as-efficient","title":"InPars-v2: Large Language Models as Efficient Dataset Generators for Information Retrieval","date":"2023-01-04","arxiv_id":"2301.01820","n_code_links":1,"syntology":null},{"paper":"/paper/multi-aspect-explainable-inductive-relation","slug":"multi-aspect-explainable-inductive-relation","title":"Multi-Aspect Explainable Inductive Relation Prediction by Sentence Transformer","date":"2023-01-04","arxiv_id":"2301.01664","n_code_links":1,"syntology":null},{"paper":null,"slug":"semi-mae-masked-autoencoders-for-semi","title":"Semi-MAE: Masked Autoencoders for Semi-supervised Vision Transformers","date":"2023-01-04","arxiv_id":"2301.01431","n_code_links":0,"syntology":null},{"paper":"/paper/spts-v2-single-point-scene-text-spotting","slug":"spts-v2-single-point-scene-text-spotting","title":"SPTS v2: Single-Point Scene Text Spotting","date":"2023-01-04","arxiv_id":"2301.01635","n_code_links":3,"syntology":{"ran":15,"of":16,"n_ran_checked":11,"n_instrument":4,"unverified":1,"pointer_only":3,"phrase":"15 ran (of which 0 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 1 violated, 10 with no contract checked; 4 where Syntology's instrument failed) · 1 unverified","official":{"repos":["bytedance/sptsv2","yuliang-liu/sptsv2","shannanyinxiang/spts"],"state":"official (archive's flag): 15 ran","n_ran":15,"n_constructed":0,"n_ran_no_instrument_failure":11,"n_unverified":1,"ran_from_kinds":["official"]}}},{"paper":"/paper/unihd-at-tsar-2022-shared-task-is-compute-all","slug":"unihd-at-tsar-2022-shared-task-is-compute-all","title":"UniHD at TSAR-2022 Shared Task: Is Compute All We Need for Lexical Simplification?","date":"2023-01-04","arxiv_id":"2301.01764","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-new-perspective-to-boost-vision-transformer","title":"A New Perspective to Boost Vision Transformer for Medical Image Classification","date":"2023-01-03","arxiv_id":"2301.00989","n_code_links":0,"syntology":null},{"paper":"/paper/cross-modal-transformer-via-coordinates","slug":"cross-modal-transformer-via-coordinates","title":"Cross Modal Transformer: Towards Fast and Robust 3D Object Detection","date":"2023-01-03","arxiv_id":"2301.01283","n_code_links":2,"syntology":null},{"paper":null,"slug":"detecting-severity-of-diabetic-retinopathy","title":"Detecting Severity of Diabetic Retinopathy from Fundus Images: A Transformer Network-based Review","date":"2023-01-03","arxiv_id":"2301.00973","n_code_links":0,"syntology":null},{"paper":"/paper/large-language-models-as-corporate-lobbyists","slug":"large-language-models-as-corporate-lobbyists","title":"Large Language Models as Corporate Lobbyists","date":"2023-01-03","arxiv_id":"2301.01181","n_code_links":1,"syntology":null},{"paper":"/paper/medical-image-segmentation-via-cascaded","slug":"medical-image-segmentation-via-cascaded","title":"Medical Image Segmentation via Cascaded Attention Decoding","date":"2023-01-03","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"modeling-the-rhythm-from-lyrics-for-melody","title":"Modeling the Rhythm from Lyrics for Melody Generation of Pop Song","date":"2023-01-03","arxiv_id":"2301.01361","n_code_links":0,"syntology":null},{"paper":null,"slug":"pie-qg-paraphrased-information-extraction-for","title":"PIE-QG: Paraphrased Information Extraction for Unsupervised Question Generation from Small Corpora","date":"2023-01-03","arxiv_id":"2301.01064","n_code_links":0,"syntology":null},{"paper":"/paper/rethinking-mobile-block-for-efficient-neural","slug":"rethinking-mobile-block-for-efficient-neural","title":"Rethinking Mobile Block for Efficient Attention-based Models","date":"2023-01-03","arxiv_id":"2301.01146","n_code_links":1,"syntology":null},{"paper":"/paper/tinymim-an-empirical-study-of-distilling-mim","slug":"tinymim-an-empirical-study-of-distilling-mim","title":"TinyMIM: An Empirical Study of Distilling MIM Pre-trained Models","date":"2023-01-03","arxiv_id":"2301.01296","n_code_links":2,"syntology":{"ran":8,"of":11,"n_ran_checked":8,"n_instrument":0,"unverified":3,"pointer_only":11,"phrase":"8 ran (of which 5 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 0 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["oliverrensu/tinymim"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":1,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/betrayed-by-captions-joint-caption-grounding","slug":"betrayed-by-captions-joint-caption-grounding","title":"Betrayed by Captions: Joint Caption Grounding and Generation for Open Vocabulary Instance Segmentation","date":"2023-01-02","arxiv_id":"2301.00805","n_code_links":2,"syntology":null},{"paper":"/paper/maud-an-expert-annotated-legal-nlp-dataset","slug":"maud-an-expert-annotated-legal-nlp-dataset","title":"MAUD: An Expert-Annotated Legal NLP Dataset for Merger Agreement Understanding","date":"2023-01-02","arxiv_id":"2301.00876","n_code_links":2,"syntology":null},{"paper":null,"slug":"multi-stage-spatio-temporal-aggregation","title":"Multi-Stage Spatio-Temporal Aggregation Transformer for Video Person Re-identification","date":"2023-01-02","arxiv_id":"2301.00531","n_code_links":0,"syntology":null},{"paper":"/paper/muse-text-to-image-generation-via-masked","slug":"muse-text-to-image-generation-via-masked","title":"Muse: Text-To-Image Generation via Masked Generative Transformers","date":"2023-01-02","arxiv_id":"2301.00704","n_code_links":5,"syntology":{"ran":19,"of":21,"n_ran_checked":8,"n_instrument":11,"unverified":2,"pointer_only":11,"phrase":"19 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 2 honoured, 5 violated, 1 with no contract checked; 11 where Syntology's instrument failed) · 2 unverified","official":null}},{"paper":null,"slug":"transformer-based-geocoding","title":"Transformer Based Geocoding","date":"2023-01-02","arxiv_id":"2301.01170","n_code_links":0,"syntology":null},{"paper":null,"slug":"understanding-political-polarisation-using","title":"Understanding Political Polarisation using Language Models: A dataset and method","date":"2023-01-02","arxiv_id":"2301.00891","n_code_links":0,"syntology":null},{"paper":"/paper/3dppe-3d-point-positional-encoding-for","slug":"3dppe-3d-point-positional-encoding-for","title":"3DPPE: 3D Point Positional Encoding for Transformer-based Multi-Camera 3D Object Detection","date":"2023-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"a-large-scale-robustness-analysis-of-video","title":"A Large-Scale Robustness Analysis of Video Action Recognition Models","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"a-simple-vision-transformer-for-weakly-semi","title":"A Simple Vision Transformer for Weakly Semi-supervised 3D Object Detection","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/adaptive-and-background-aware-vision","slug":"adaptive-and-background-aware-vision","title":"Adaptive and Background-Aware Vision Transformer for Real-Time UAV Tracking","date":"2023-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/adversarial-normalization-i-can-visualize","slug":"adversarial-normalization-i-can-visualize","title":"Adversarial Normalization: I Can Visualize Everything (ICE)","date":"2023-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"attentionshift-iteratively-estimated-part","title":"AttentionShift: Iteratively Estimated Part-Based Attention Map for Pointly Supervised Instance Segmentation","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"aunet-learning-relations-between-action-units","title":"AUNet: Learning Relations Between Action Units for Face Forgery Detection","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-knowledge-distillation-via-monte","title":"Automated Knowledge Distillation via Monte Carlo Tree Search","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bit-shrinking-limiting-instantaneous","title":"Bit-Shrinking: Limiting Instantaneous Sharpness for Improving Post-Training Quantization","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"boosting-whole-slide-image-classification","title":"Boosting Whole Slide Image Classification from the Perspectives of Distribution, Correlation and Magnification","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"building-vision-transformers-with-hierarchy","title":"Building Vision Transformers with Hierarchy Aware Feature Aggregation","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"bus-efficient-and-effective-vision-language-1","title":"BUS: Efficient and Effective Vision-Language Pre-Training with Bottom-Up Patch Summarization.","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/chasing-clouds-differentiable-volumetric","slug":"chasing-clouds-differentiable-volumetric","title":"Chasing Clouds: Differentiable Volumetric Rasterisation of Point Clouds as a Highly Efficient and Accurate Loss for Large-Scale Deformable 3D Registration","date":"2023-01-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"clipping-distilling-clip-based-models-with-a","title":"CLIPPING: Distilling CLIP-Based Models With a Student Base for Video-Language Retrieval","date":"2023-01-01","arxiv_id":null,"n_code_links":0,"syntology":null}],"record_sha256":"c0c24bf46b0d633510816706ddf9e8aeb871376cbc17e3d2e55c69f92781c69f","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}