{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/dropout/papers/152","list_of":"/method/dropout","method":"Dropout","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":152,"pages_in_order":275,"rows_per_page":100,"rows":[15101,15200],"of":27472,"counts":{"archive_papers_tagged":27472,"with_a_code_link":12129,"where_syntology_ran_a_sample":3620,"not_listed_spam_title":0,"listed":27472,"listed_where_code_ran":3620,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3044,"every_run_a_failure_of_syntologys_instrument":576,"listed_with_a_run_with_no_instrument_failure":3044,"listed_every_run_a_failure_of_syntologys_instrument":576,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/dropout","prev":"/method/dropout/papers/151","next":"/method/dropout/papers/153","papers":[{"paper":"/paper/exaranker-explanation-augmented-neural-ranker","slug":"exaranker-explanation-augmented-neural-ranker","title":"ExaRanker: Explanation-Augmented Neural Ranker","date":"2023-01-25","arxiv_id":"2301.10521","n_code_links":1,"syntology":null},{"paper":null,"slug":"qualitative-analysis-of-a-graph-transformer","title":"Qualitative Analysis of a Graph Transformer Approach to Addressing Hate Speech: Adapting to Dynamically Changing Content","date":"2023-01-25","arxiv_id":"2301.10871","n_code_links":0,"syntology":null},{"paper":"/paper/transfer-learning-in-deep-learning-models-for","slug":"transfer-learning-in-deep-learning-models-for","title":"Transfer Learning in Deep Learning Models for Building Load Forecasting: Case of Limited Data","date":"2023-01-25","arxiv_id":"2301.10663","n_code_links":2,"syntology":null},{"paper":"/paper/videberta-a-powerful-pre-trained-language","slug":"videberta-a-powerful-pre-trained-language","title":"ViDeBERTa: A powerful pre-trained language model for Vietnamese","date":"2023-01-25","arxiv_id":"2301.10439","n_code_links":1,"syntology":{"ran":2,"of":4,"n_ran_checked":2,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["hysonlab/videberta"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"a-stability-analysis-of-fine-tuning-a-pre","title":"A Stability Analysis of Fine-Tuning a Pre-Trained Model","date":"2023-01-24","arxiv_id":"2301.09820","n_code_links":0,"syntology":null},{"paper":"/paper/a-watermark-for-large-language-models","slug":"a-watermark-for-large-language-models","title":"A Watermark for Large Language Models","date":"2023-01-24","arxiv_id":"2301.10226","n_code_links":8,"syntology":{"ran":11,"of":16,"n_ran_checked":10,"n_instrument":1,"unverified":5,"pointer_only":5,"phrase":"11 ran (of which 7 constructed an object rather than computing a result; 10 with no instrument failure: 1 honoured, 1 violated, 8 with no contract checked; 1 where Syntology's instrument failed) · 5 unverified","official":{"repos":["jwkirchenbauer/lm-watermarking"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/audience-centric-natural-language-generation","slug":"audience-centric-natural-language-generation","title":"Audience-Centric Natural Language Generation via Style Infusion","date":"2023-01-24","arxiv_id":"2301.10283","n_code_links":1,"syntology":null},{"paper":null,"slug":"can-very-large-pretrained-language-models","title":"The Next Chapter: A Study of Large Language Models in Storytelling","date":"2023-01-24","arxiv_id":"2301.09790","n_code_links":0,"syntology":null},{"paper":"/paper/climax-a-foundation-model-for-weather-and","slug":"climax-a-foundation-model-for-weather-and","title":"ClimaX: A foundation model for weather and climate","date":"2023-01-24","arxiv_id":"2301.10343","n_code_links":1,"syntology":null},{"paper":null,"slug":"large-language-models-as-fiduciaries-a-case","title":"Large Language Models as Fiduciaries: A Case Study Toward Robustly Communicating With Artificial Intelligence Through Legal Standards","date":"2023-01-24","arxiv_id":"2301.10095","n_code_links":0,"syntology":null},{"paper":null,"slug":"large-language-models-can-segment-narrative","title":"Large language models can segment narrative events similarly to humans","date":"2023-01-24","arxiv_id":"2301.10297","n_code_links":0,"syntology":null},{"paper":"/paper/model-soups-to-increase-inference-without","slug":"model-soups-to-increase-inference-without","title":"Model soups to increase inference without increasing compute time","date":"2023-01-24","arxiv_id":"2301.10092","n_code_links":1,"syntology":null},{"paper":null,"slug":"multitask-instruction-based-prompting-for","title":"Multitask Instruction-based Prompting for Fallacy Recognition","date":"2023-01-24","arxiv_id":"2301.09992","n_code_links":0,"syntology":null},{"paper":null,"slug":"smart-self-supervised-multi-task-pretraining","title":"SMART: Self-supervised Multi-task pretrAining with contRol Transformers","date":"2023-01-24","arxiv_id":"2301.09816","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-simple-recipe-for-competitive-low-compute","title":"A Simple Recipe for Competitive Low-compute Self supervised Vision Models","date":"2023-01-23","arxiv_id":"2301.09451","n_code_links":0,"syntology":null},{"paper":null,"slug":"ai-model-gpt-3-dis-informs-us-better-than","title":"AI model GPT-3 (dis)informs us better than humans","date":"2023-01-23","arxiv_id":"2301.11924","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-mental-health-dialogue-system","title":"Deep Learning Mental Health Dialogue System","date":"2023-01-23","arxiv_id":"2301.09412","n_code_links":0,"syntology":null},{"paper":"/paper/efficient-language-model-training-through","slug":"efficient-language-model-training-through","title":"Efficient Language Model Training through Cross-Lingual and Progressive Transfer Learning","date":"2023-01-23","arxiv_id":"2301.09626","n_code_links":1,"syntology":null},{"paper":"/paper/fully-transformer-based-biomarker-prediction","slug":"fully-transformer-based-biomarker-prediction","title":"Fully transformer-based biomarker prediction from colorectal cancer histology: a large-scale multicentric study","date":"2023-01-23","arxiv_id":"2301.09617","n_code_links":2,"syntology":null},{"paper":"/paper/injecting-the-bm25-score-as-text-improves","slug":"injecting-the-bm25-score-as-text-improves","title":"Injecting the BM25 Score as Text Improves BERT-Based Re-rankers","date":"2023-01-23","arxiv_id":"2301.09728","n_code_links":1,"syntology":null},{"paper":"/paper/istvt-interpretable-spatial-temporal-video","slug":"istvt-interpretable-spatial-temporal-video","title":"ISTVT: Interpretable Spatial-Temporal Video Transformer for Deepfake Detection","date":"2023-01-23","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"learning-to-view-decision-transformers-for","title":"Learning to View: Decision Transformers for Active Object Detection","date":"2023-01-23","arxiv_id":"2301.09544","n_code_links":0,"syntology":null},{"paper":"/paper/local-window-attention-transformer-for","slug":"local-window-attention-transformer-for","title":"Local Window Attention Transformer for Polarimetric SAR Image Classification","date":"2023-01-23","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/stockemotions-discover-investor-emotions-for","slug":"stockemotions-discover-investor-emotions-for","title":"StockEmotions: Discover Investor Emotions for Financial Sentiment Analysis and Multivariate Time Series","date":"2023-01-23","arxiv_id":"2301.09279","n_code_links":2,"syntology":null},{"paper":null,"slug":"apples-and-oranges-assessing-image-quality","title":"Apples and Oranges? Assessing Image Quality over Content Recognition","date":"2023-01-22","arxiv_id":"2301.09190","n_code_links":0,"syntology":null},{"paper":"/paper/debiasing-the-cloze-task-in-sequential","slug":"debiasing-the-cloze-task-in-sequential","title":"Debiasing the Cloze Task in Sequential Recommendation with Bidirectional Transformers","date":"2023-01-22","arxiv_id":"2301.09210","n_code_links":1,"syntology":null},{"paper":null,"slug":"estimation-of-sea-state-parameters-from-ship","title":"Estimation of Sea State Parameters from Ship Motion Responses Using Attention-based Neural Networks","date":"2023-01-21","arxiv_id":"2301.08949","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-methods-for-building-dialects","slug":"exploring-methods-for-building-dialects","title":"Exploring Methods for Building Dialects-Mandarin Code-Mixing Corpora: A Case Study in Taiwanese Hokkien","date":"2023-01-21","arxiv_id":"2301.08937","n_code_links":1,"syntology":null},{"paper":null,"slug":"slice-transformer-and-self-supervised","title":"Slice Transformer and Self-supervised Learning for 6DoF Localization in 3D Point Cloud Maps","date":"2023-01-21","arxiv_id":"2301.08957","n_code_links":0,"syntology":null},{"paper":null,"slug":"stress-test-for-bert-and-deep-models","title":"Stress Test for BERT and Deep Models: Predicting Words from Italian Poetry","date":"2023-01-21","arxiv_id":"2302.09303","n_code_links":0,"syntology":null},{"paper":null,"slug":"superscaler-supporting-flexible-dnn","title":"SuperScaler: Supporting Flexible DNN Parallelization via a Unified Abstraction","date":"2023-01-21","arxiv_id":"2301.08984","n_code_links":0,"syntology":null},{"paper":"/paper/time-conditioned-generative-modeling-of","slug":"time-conditioned-generative-modeling-of","title":"Time-Conditioned Generative Modeling of Object-Centric Representations for Video Decomposition and Prediction","date":"2023-01-21","arxiv_id":"2301.08951","n_code_links":1,"syntology":null},{"paper":null,"slug":"accelerating-multi-agent-planning-using-graph","title":"Accelerating Multi-Agent Planning Using Graph Transformers with Bounded Suboptimality","date":"2023-01-20","arxiv_id":"2301.08451","n_code_links":0,"syntology":null},{"paper":"/paper/is-chatgpt-a-good-translator-a-preliminary","slug":"is-chatgpt-a-good-translator-a-preliminary","title":"Is ChatGPT A Good Translator? Yes With GPT-4 As The Engine","date":"2023-01-20","arxiv_id":"2301.08745","n_code_links":1,"syntology":null},{"paper":null,"slug":"ontology-pre-training-for-poison-prediction","title":"Ontology Pre-training for Poison Prediction","date":"2023-01-20","arxiv_id":"2301.08577","n_code_links":0,"syntology":null},{"paper":"/paper/phoneme-level-bert-for-enhanced-prosody-of","slug":"phoneme-level-bert-for-enhanced-prosody-of","title":"Phoneme-Level BERT for Enhanced Prosody of Text-to-Speech with Grapheme Predictions","date":"2023-01-20","arxiv_id":"2301.08810","n_code_links":2,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":null}},{"paper":null,"slug":"which-features-are-learned-by-codebert-an","title":"Which Features are Learned by CodeBert: An Empirical Study of the BERT-based Source Code Representation Learning","date":"2023-01-20","arxiv_id":"2301.08427","n_code_links":0,"syntology":null},{"paper":"/paper/batch-prompting-efficient-inference-with","slug":"batch-prompting-efficient-inference-with","title":"Batch Prompting: Efficient Inference with Large Language Model APIs","date":"2023-01-19","arxiv_id":"2301.08721","n_code_links":2,"syntology":null},{"paper":"/paper/diagnose-like-a-pathologist-transformer","slug":"diagnose-like-a-pathologist-transformer","title":"Diagnose Like a Pathologist: Transformer-Enabled Hierarchical Attention-Guided Multiple Instance Learning for Whole Slide Image Classification","date":"2023-01-19","arxiv_id":"2301.08125","n_code_links":1,"syntology":null},{"paper":null,"slug":"fe-tcm-filter-enhanced-transformer-click","title":"FE-TCM: Filter-Enhanced Transformer Click Model for Web Search","date":"2023-01-19","arxiv_id":"2301.07854","n_code_links":0,"syntology":null},{"paper":"/paper/m3e-yolo-a-new-lightweight-network-for","slug":"m3e-yolo-a-new-lightweight-network-for","title":"M3E-Yolo: A New Lightweight Network for Traffic Sign Recognition","date":"2023-01-19","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/medsegdiff-v2-diffusion-based-medical-image","slug":"medsegdiff-v2-diffusion-based-medical-image","title":"MedSegDiff-V2: Diffusion based Medical Image Segmentation with Transformer","date":"2023-01-19","arxiv_id":"2301.11798","n_code_links":2,"syntology":{"ran":17,"of":21,"n_ran_checked":12,"n_instrument":5,"unverified":4,"pointer_only":10,"phrase":"17 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 2 honoured, 0 violated, 10 with no contract checked; 5 where Syntology's instrument failed) · 4 unverified","official":{"repos":["kidswithtokens/medsegdiff","wujunde/medsegdiff"],"state":"official (archive's flag): 17 ran","n_ran":17,"n_constructed":0,"n_ran_no_instrument_failure":12,"n_unverified":4,"ran_from_kinds":["official"]}}},{"paper":"/paper/graphix-t5-mixing-pre-trained-transformers","slug":"graphix-t5-mixing-pre-trained-transformers","title":"Graphix-T5: Mixing Pre-Trained Transformers with Graph-Aware Layers for Text-to-SQL Parsing","date":"2023-01-18","arxiv_id":"2301.07507","n_code_links":1,"syntology":null},{"paper":null,"slug":"cooperation-learning-enhanced-colonic-polyp","title":"Cooperation Learning Enhanced Colonic Polyp Segmentation Based on Transformer-CNN Fusion","date":"2023-01-17","arxiv_id":"2301.06892","n_code_links":0,"syntology":null},{"paper":"/paper/sat-size-aware-transformer-for-3d-point-cloud","slug":"sat-size-aware-transformer-for-3d-point-cloud","title":"SAT: Size-Aware Transformer for 3D Point Cloud Semantic Segmentation","date":"2023-01-17","arxiv_id":"2301.06869","n_code_links":0,"syntology":null},{"paper":"/paper/swindepth-unsupervised-depth-estimation-using","slug":"swindepth-unsupervised-depth-estimation-using","title":"SwinDepth: Unsupervised Depth Estimation using Monocular Sequences via Swin Transformer and Densely Cascaded Network","date":"2023-01-17","arxiv_id":"2301.06715","n_code_links":1,"syntology":null},{"paper":"/paper/tracing-and-manipulating-intermediate-values","slug":"tracing-and-manipulating-intermediate-values","title":"Tracing and Manipulating Intermediate Values in Neural Math Problem Solvers","date":"2023-01-17","arxiv_id":"2301.06758","n_code_links":1,"syntology":null},{"paper":null,"slug":"transformer-based-implementation-for","title":"Transformer Based Implementation for Automatic Book Summarization","date":"2023-01-17","arxiv_id":"2301.07057","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-2-uav-application-aware-content-and-network","title":"A$^2$-UAV: Application-Aware Content and Network Optimization of Edge-Assisted UAV Systems","date":"2023-01-16","arxiv_id":"2301.06363","n_code_links":0,"syntology":null},{"paper":"/paper/an-error-guided-correction-model-for-chinese","slug":"an-error-guided-correction-model-for-chinese","title":"An Error-Guided Correction Model for Chinese Spelling Error Correction","date":"2023-01-16","arxiv_id":"2301.06323","n_code_links":1,"syntology":null},{"paper":null,"slug":"bayesspeech-a-bayesian-transformer-network","title":"BayesSpeech: A Bayesian Transformer Network for Automatic Speech Recognition","date":"2023-01-16","arxiv_id":"2301.11276","n_code_links":0,"syntology":null},{"paper":null,"slug":"masked-vector-quantization","title":"Masked Vector Quantization","date":"2023-01-16","arxiv_id":"2301.06626","n_code_links":0,"syntology":null},{"paper":"/paper/tdstf-transformer-based-diffusion","slug":"tdstf-transformer-based-diffusion","title":"A Transformer-based Diffusion Probabilistic Model for Heart Rate and Blood Pressure Forecasting in Intensive Care Unit","date":"2023-01-16","arxiv_id":"2301.06625","n_code_links":1,"syntology":null},{"paper":"/paper/tedb-system-description-to-a-shared-task-on","slug":"tedb-system-description-to-a-shared-task-on","title":"TEDB System Description to a Shared Task on Euphemism Detection 2022","date":"2023-01-16","arxiv_id":"2301.06602","n_code_links":1,"syntology":null},{"paper":"/paper/dsvt-dynamic-sparse-voxel-transformer-with","slug":"dsvt-dynamic-sparse-voxel-transformer-with","title":"DSVT: Dynamic Sparse Voxel Transformer with Rotated Sets","date":"2023-01-15","arxiv_id":"2301.06051","n_code_links":4,"syntology":null},{"paper":null,"slug":"improving-noise-robustness-for-spoken-content","title":"Improving Noise Robustness for Spoken Content Retrieval using Semi-supervised ASR and N-best Transcripts for BERT-based Ranking Models","date":"2023-01-15","arxiv_id":"2301.06056","n_code_links":0,"syntology":null},{"paper":"/paper/t2m-gpt-generating-human-motion-from-textual","slug":"t2m-gpt-generating-human-motion-from-textual","title":"T2M-GPT: Generating Human Motion from Textual Descriptions with Discrete Representations","date":"2023-01-15","arxiv_id":"2301.06052","n_code_links":1,"syntology":{"ran":5,"of":7,"n_ran_checked":5,"n_instrument":0,"unverified":2,"pointer_only":0,"phrase":"5 ran (of which 4 constructed an object rather than computing a result; 5 with no instrument failure: 1 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified","official":{"repos":["Mael-zys/T2M-GPT"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":4,"n_ran_no_instrument_failure":5,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"salient-sign-detection-in-safe-autonomous","title":"Salient Sign Detection In Safe Autonomous Driving: AI Which Reasons Over Full Visual Context","date":"2023-01-14","arxiv_id":"2301.05804","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-speech-and-text-based","title":"Automated speech- and text-based classification of neuropsychiatric conditions in a multidiagnostic setting","date":"2023-01-13","arxiv_id":"2301.06916","n_code_links":0,"syntology":null},{"paper":null,"slug":"text-to-point-cloud-localization-with","title":"Text to Point Cloud Localization with Relation-Enhanced Transformer","date":"2023-01-13","arxiv_id":"2301.05372","n_code_links":0,"syntology":null},{"paper":"/paper/adversarial-adaptation-for-french-named","slug":"adversarial-adaptation-for-french-named","title":"Adversarial Adaptation for French Named Entity Recognition","date":"2023-01-12","arxiv_id":"2301.05220","n_code_links":1,"syntology":null},{"paper":null,"slug":"sacdnet-towards-early-type-2-diabetes","title":"SACDNet: Towards Early Type 2 Diabetes Prediction with Uncertainty for Electronic Health Records","date":"2023-01-12","arxiv_id":"2301.04844","n_code_links":0,"syntology":null},{"paper":"/paper/vits-for-sits-vision-transformers-for","slug":"vits-for-sits-vision-transformers-for","title":"ViTs for SITS: Vision Transformers for Satellite Image Time Series","date":"2023-01-12","arxiv_id":"2301.04944","n_code_links":3,"syntology":{"ran":11,"of":11,"n_ran_checked":11,"n_instrument":0,"unverified":0,"pointer_only":0,"phrase":"11 ran (of which 3 constructed an object rather than computing a result; 11 with no instrument failure: 0 honoured, 0 violated, 11 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","official":{"repos":["michaeltrs/deepsatmodels"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":3,"n_ran_no_instrument_failure":3,"n_unverified":0,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/adapointr-diverse-point-cloud-completion-with","slug":"adapointr-diverse-point-cloud-completion-with","title":"AdaPoinTr: Diverse Point Cloud Completion with Adaptive Geometry-Aware Transformers","date":"2023-01-11","arxiv_id":"2301.04545","n_code_links":1,"syntology":{"ran":4,"of":6,"n_ran_checked":2,"n_instrument":2,"unverified":2,"pointer_only":3,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 2 unverified","official":{"repos":["yuxumin/PoinTr"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"anomalies-representations-and-self","title":"Anomalies, Representations, and Self-Supervision","date":"2023-01-11","arxiv_id":"2301.04660","n_code_links":0,"syntology":null},{"paper":null,"slug":"dynamic-background-reconstruction-via","title":"Dynamic Background Reconstruction via MAE for Infrared Small Target Detection","date":"2023-01-11","arxiv_id":"2301.04497","n_code_links":0,"syntology":null},{"paper":"/paper/gpt-as-knowledge-worker-a-zero-shot","slug":"gpt-as-knowledge-worker-a-zero-shot","title":"GPT as Knowledge Worker: A Zero-Shot Evaluation of (AI)CPA Capabilities","date":"2023-01-11","arxiv_id":"2301.04408","n_code_links":1,"syntology":null},{"paper":"/paper/head-free-lightweight-semantic-segmentation","slug":"head-free-lightweight-semantic-segmentation","title":"Head-Free Lightweight Semantic Segmentation with Linear Transformer","date":"2023-01-11","arxiv_id":"2301.04648","n_code_links":1,"syntology":null},{"paper":"/paper/narrowbert-accelerating-masked-language-model","slug":"narrowbert-accelerating-masked-language-model","title":"NarrowBERT: Accelerating Masked Language Model Pretraining and Inference","date":"2023-01-11","arxiv_id":"2301.04761","n_code_links":1,"syntology":{"ran":0,"of":2,"n_ran_checked":0,"n_instrument":0,"unverified":2,"pointer_only":2,"phrase":"0 ran · 2 unverified","official":{"repos":["lihaoxin2020/narrowbert"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":2,"ran_from_kinds":[]}}},{"paper":null,"slug":"topics-in-contextualised-attention-embeddings","title":"Topics in Contextualised Attention Embeddings","date":"2023-01-11","arxiv_id":"2301.04339","n_code_links":0,"syntology":null},{"paper":null,"slug":"does-image-resolution-impact-chest-x-ray","title":"Does image resolution impact chest X-ray based fine-grained Tuberculosis-consistent lesion segmentation?","date":"2023-01-10","arxiv_id":"2301.04032","n_code_links":0,"syntology":null},{"paper":"/paper/dropcov-a-simple-yet-effective-method-for","slug":"dropcov-a-simple-yet-effective-method-for","title":"DropCov: A Simple yet Effective Method for Improving Deep Architectures","date":"2023-01-10","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"language-models-sounds-the-death-knell-of","title":"Language Models sounds the Death Knell of Knowledge Graphs","date":"2023-01-10","arxiv_id":"2301.03980","n_code_links":0,"syntology":null},{"paper":"/paper/predicting-hateful-discussions-on-reddit","slug":"predicting-hateful-discussions-on-reddit","title":"Predicting Hateful Discussions on Reddit using Graph Transformer Networks and Communal Context","date":"2023-01-10","arxiv_id":"2301.04248","n_code_links":1,"syntology":null},{"paper":null,"slug":"recommending-root-cause-and-mitigation-steps","title":"Recommending Root-Cause and Mitigation Steps for Cloud Incidents using Large Language Models","date":"2023-01-10","arxiv_id":"2301.03797","n_code_links":0,"syntology":null},{"paper":null,"slug":"streaming-punctuation-a-novel-punctuation","title":"Streaming Punctuation: A Novel Punctuation Technique Leveraging Bidirectional Context for Continuous Speech Recognition","date":"2023-01-10","arxiv_id":"2301.03819","n_code_links":0,"syntology":null},{"paper":"/paper/there-is-no-big-brother-or-small-brother","slug":"there-is-no-big-brother-or-small-brother","title":"There is No Big Brother or Small Brother: Knowledge Infusion in Language Models for Link Prediction and Question Answering","date":"2023-01-10","arxiv_id":"2301.04013","n_code_links":2,"syntology":null},{"paper":"/paper/unsupervised-mandarin-cantonese-machine","slug":"unsupervised-mandarin-cantonese-machine","title":"Unsupervised Mandarin-Cantonese Machine Translation","date":"2023-01-10","arxiv_id":"2301.03971","n_code_links":1,"syntology":null},{"paper":"/paper/a-study-on-the-generality-of-neural-network","slug":"a-study-on-the-generality-of-neural-network","title":"A Study on the Generality of Neural Network Structures for Monocular Depth Estimation","date":"2023-01-09","arxiv_id":"2301.03169","n_code_links":1,"syntology":null},{"paper":"/paper/advances-in-medical-image-analysis-with","slug":"advances-in-medical-image-analysis-with","title":"Advances in Medical Image Analysis with Vision Transformers: A Comprehensive Review","date":"2023-01-09","arxiv_id":"2301.03505","n_code_links":1,"syntology":null},{"paper":"/paper/an-impartial-transformer-for-story","slug":"an-impartial-transformer-for-story","title":"An Impartial Transformer for Story Visualization","date":"2023-01-09","arxiv_id":"2301.03563","n_code_links":0,"syntology":null},{"paper":"/paper/demt-deformable-mixer-transformer-for-multi","slug":"demt-deformable-mixer-transformer-for-multi","title":"DeMT: Deformable Mixer Transformer for Multi-Task Learning of Dense Prediction","date":"2023-01-09","arxiv_id":"2301.03461","n_code_links":1,"syntology":{"ran":3,"of":6,"n_ran_checked":3,"n_instrument":0,"unverified":3,"pointer_only":6,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 3 with no instrument failure: 0 honoured, 0 violated, 3 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["yangyangxu0/demt"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":3,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":"/paper/designing-bert-for-convolutional-networks","slug":"designing-bert-for-convolutional-networks","title":"Designing BERT for Convolutional Networks: Sparse and Hierarchical Masked Modeling","date":"2023-01-09","arxiv_id":"2301.03580","n_code_links":2,"syntology":{"ran":13,"of":14,"n_ran_checked":12,"n_instrument":1,"unverified":1,"pointer_only":2,"phrase":"13 ran (of which 0 constructed an object rather than computing a result; 12 with no instrument failure: 0 honoured, 0 violated, 12 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["keyu-tian/spark"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":null,"slug":"logically-at-factify-2023-a-multi-modal-fact","title":"Logically at Factify 2: A Multi-Modal Fact Checking System Based on Evidence Retrieval techniques and Transformer Encoder Architecture","date":"2023-01-09","arxiv_id":"2301.03127","n_code_links":0,"syntology":null},{"paper":null,"slug":"online-fake-review-detection-using-supervised","title":"Online Fake Review Detection Using Supervised Machine Learning And BERT Model","date":"2023-01-09","arxiv_id":"2301.03225","n_code_links":0,"syntology":null},{"paper":null,"slug":"universal-multimodal-representation-for","title":"Universal Multimodal Representation for Language Understanding","date":"2023-01-09","arxiv_id":"2301.03344","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-generation-of-german-drama-texts","title":"Automatic Generation of German Drama Texts Using Fine Tuned GPT-2 Models","date":"2023-01-08","arxiv_id":"2301.03119","n_code_links":0,"syntology":null},{"paper":"/paper/deepmatcher-a-deep-transformer-based-network","slug":"deepmatcher-a-deep-transformer-based-network","title":"DeepMatcher: A Deep Transformer-based Network for Robust and Accurate Local Feature Matching","date":"2023-01-08","arxiv_id":"2301.02993","n_code_links":1,"syntology":null},{"paper":"/paper/hrtransnet-hrformer-driven-two-modality","slug":"hrtransnet-hrformer-driven-two-modality","title":"HRTransNet: HRFormer-Driven Two-Modality Salient Object Detection","date":"2023-01-08","arxiv_id":"2301.03036","n_code_links":1,"syntology":null},{"paper":"/paper/app-review-driven-collaborative-bug-finding","slug":"app-review-driven-collaborative-bug-finding","title":"App Review Driven Collaborative Bug Finding","date":"2023-01-07","arxiv_id":"2301.02818","n_code_links":1,"syntology":null},{"paper":"/paper/rlas-biabc-a-reinforcement-learning-based","slug":"rlas-biabc-a-reinforcement-learning-based","title":"RLAS-BIABC: A Reinforcement Learning-Based Answer Selection Using the BERT Model Boosted by an Improved ABC Algorithm","date":"2023-01-07","arxiv_id":"2301.02807","n_code_links":0,"syntology":null},{"paper":"/paper/codetalker-speech-driven-3d-facial-animation","slug":"codetalker-speech-driven-3d-facial-animation","title":"CodeTalker: Speech-Driven 3D Facial Animation with Discrete Motion Prior","date":"2023-01-06","arxiv_id":"2301.02379","n_code_links":1,"syntology":{"ran":8,"of":13,"n_ran_checked":8,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 0 honoured, 1 violated, 7 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","official":{"repos":["Doubiiu/CodeTalker"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":5,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"conditional-generation-of-paired-antibody","title":"Generative Antibody Design for Complementary Chain Pairing Sequences through Encoder-Decoder Language Model","date":"2023-01-06","arxiv_id":"2301.02748","n_code_links":0,"syntology":null},{"paper":null,"slug":"deep-learning-models-in-medical-image","title":"Deep-learning models in medical image analysis: Detection of esophagitis from the Kvasir Dataset","date":"2023-01-06","arxiv_id":"2301.02390","n_code_links":0,"syntology":null},{"paper":null,"slug":"does-compressing-activations-help-model","title":"Does compressing activations help model parallel training?","date":"2023-01-06","arxiv_id":"2301.02654","n_code_links":0,"syntology":null},{"paper":"/paper/exploring-efficient-few-shot-adaptation-for","slug":"exploring-efficient-few-shot-adaptation-for","title":"Exploring Efficient Few-shot Adaptation for Vision Transformers","date":"2023-01-06","arxiv_id":"2301.02419","n_code_links":1,"syntology":null},{"paper":null,"slug":"multi-genre-music-transformer-composing-full","title":"Multi-Genre Music Transformer -- Composing Full Length Musical Piece","date":"2023-01-06","arxiv_id":"2301.02385","n_code_links":0,"syntology":null},{"paper":null,"slug":"systems-for-parallel-and-distributed-large","title":"Systems for Parallel and Distributed Large-Model Deep Learning Training","date":"2023-01-06","arxiv_id":"2301.02691","n_code_links":0,"syntology":null},{"paper":null,"slug":"task-aware-feature-extraction-framework-for","title":"Adaptive Pattern Extraction Multi-Task Learning for Multi-Step Conversion Estimations","date":"2023-01-06","arxiv_id":"2301.02494","n_code_links":0,"syntology":null},{"paper":"/paper/adaptively-clustering-neighbor-elements-for","slug":"adaptively-clustering-neighbor-elements-for","title":"Adaptively Clustering Neighbor Elements for Image-Text Generation","date":"2023-01-05","arxiv_id":"2301.01955","n_code_links":1,"syntology":null}],"record_sha256":"8f0a79be55c82413e8813388dcb70456047df8ab755b9d1cf881daaf03df3b0a","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}