{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/softmax/papers/319","list_of":"/method/softmax","method":"Softmax","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":319,"pages_in_order":375,"rows_per_page":100,"rows":[31801,31900],"of":37443,"counts":{"archive_papers_tagged":37443,"with_a_code_link":15869,"where_syntology_ran_a_sample":4578,"not_listed_spam_title":0,"listed":37443,"listed_where_code_ran":4578,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":3835,"every_run_a_failure_of_syntologys_instrument":743,"listed_with_a_run_with_no_instrument_failure":3835,"listed_every_run_a_failure_of_syntologys_instrument":743,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/softmax","prev":"/method/softmax/papers/318","next":"/method/softmax/papers/320","papers":[{"paper":null,"slug":"self-supervised-learning-with-cross-modal","title":"Self-Supervised learning with cross-modal transformers for emotion recognition","date":"2020-11-20","arxiv_id":"2011.10652","n_code_links":0,"syntology":null},{"paper":"/paper/dct-mask-discrete-cosine-transform-mask","slug":"dct-mask-discrete-cosine-transform-mask","title":"DCT-Mask: Discrete Cosine Transform Mask Representation for Instance Segmentation","date":"2020-11-19","arxiv_id":"2011.09876","n_code_links":1,"syntology":null},{"paper":"/paper/fact-level-extractive-summarization-with","slug":"fact-level-extractive-summarization-with","title":"Fact-level Extractive Summarization with Hierarchical Graph Mask on BERT","date":"2020-11-19","arxiv_id":"2011.09739","n_code_links":1,"syntology":null},{"paper":null,"slug":"latent-separated-global-prediction-for","title":"Causal Contextual Prediction for Learned Image Compression","date":"2020-11-19","arxiv_id":"2011.09704","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-predict-the-3d-layout-of-a-scene","title":"Learning to Predict the 3D Layout of a Scene","date":"2020-11-19","arxiv_id":"2011.09977","n_code_links":0,"syntology":null},{"paper":null,"slug":"persuasive-dialogue-understanding-the","title":"Persuasive Dialogue Understanding: the Baselines and Negative Results","date":"2020-11-19","arxiv_id":"2011.09954","n_code_links":0,"syntology":null},{"paper":null,"slug":"reassert-deep-learning-for-assert-generation","title":"ReAssert: Deep Learning for Assert Generation","date":"2020-11-19","arxiv_id":"2011.09784","n_code_links":0,"syntology":null},{"paper":null,"slug":"unifying-instance-and-panoptic-segmentation","title":"Unifying Instance and Panoptic Segmentation with Dynamic Rank-1 Convolutions","date":"2020-11-19","arxiv_id":"2011.09796","n_code_links":0,"syntology":null},{"paper":"/paper/attentivenas-improving-neural-architecture","slug":"attentivenas-improving-neural-architecture","title":"AttentiveNAS: Improving Neural Architecture Search via Attentive Sampling","date":"2020-11-18","arxiv_id":"2011.09011","n_code_links":2,"syntology":null},{"paper":null,"slug":"diverse-and-non-redundant-answer-set","title":"Diverse and Non-redundant Answer Set Extraction on Community QA based on DPPs","date":"2020-11-18","arxiv_id":"2011.09140","n_code_links":0,"syntology":null},{"paper":null,"slug":"do-fine-tuned-commonsense-language-models","title":"Do Fine-tuned Commonsense Language Models Really Generalize?","date":"2020-11-18","arxiv_id":"2011.09159","n_code_links":0,"syntology":null},{"paper":"/paper/end-to-end-object-detection-with-adaptive","slug":"end-to-end-object-detection-with-adaptive","title":"End-to-End Object Detection with Adaptive Clustering Transformer","date":"2020-11-18","arxiv_id":"2011.09315","n_code_links":1,"syntology":null},{"paper":"/paper/kidney-level-lupus-nephritis-classification","slug":"kidney-level-lupus-nephritis-classification","title":"Kidney Level Lupus Nephritis Classification using Uncertainty Guided Bayesian Convolutional Neural Networks","date":"2020-11-18","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"layer-wise-data-free-cnn-compression","title":"Layer-Wise Data-Free CNN Compression","date":"2020-11-18","arxiv_id":"2011.09058","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-control-for-transmission-and","title":"Learning control for transmission and navigation with a mobile robot under unknown communication rates","date":"2020-11-18","arxiv_id":"2011.09193","n_code_links":0,"syntology":null},{"paper":null,"slug":"palomino-ochoa-at-semeval-2020-task-9-robust","title":"Palomino-Ochoa at SemEval-2020 Task 9: Robust System based on Transformer for Code-Mixed Sentiment Classification","date":"2020-11-18","arxiv_id":"2011.09448","n_code_links":0,"syntology":null},{"paper":null,"slug":"proposing-method-to-increase-the-detection","title":"Proposing method to Increase the detection accuracy of stomach cancer based on colour and lint features of tongue using CNN and SVM","date":"2020-11-18","arxiv_id":"2011.09962","n_code_links":0,"syntology":null},{"paper":"/paper/sequence-level-mixed-sample-data-augmentation","slug":"sequence-level-mixed-sample-data-augmentation","title":"Sequence-Level Mixed Sample Data Augmentation","date":"2020-11-18","arxiv_id":"2011.09039","n_code_links":1,"syntology":null},{"paper":null,"slug":"shaping-deep-feature-space-towards-gaussian","title":"Shaping Deep Feature Space towards Gaussian Mixture for Visual Classification","date":"2020-11-18","arxiv_id":"2011.09066","n_code_links":0,"syntology":null},{"paper":null,"slug":"the-ubiqus-english-inuktitut-system-for-wmt20","title":"The Ubiqus English-Inuktitut System for WMT20","date":"2020-11-18","arxiv_id":"2011.09249","n_code_links":0,"syntology":null},{"paper":null,"slug":"tie-your-embeddings-down-cross-modal-latent","title":"Tie Your Embeddings Down: Cross-Modal Latent Spaces for End-to-end Spoken Language Understanding","date":"2020-11-18","arxiv_id":"2011.09044","n_code_links":0,"syntology":null},{"paper":"/paper/up-detr-unsupervised-pre-training-for-object","slug":"up-detr-unsupervised-pre-training-for-object","title":"UP-DETR: Unsupervised Pre-training for Object Detection with Transformers","date":"2020-11-18","arxiv_id":"2011.09094","n_code_links":2,"syntology":null},{"paper":null,"slug":"a-comparative-approach-on-detecting-multi","title":"A comparative approach on detecting multi-lingual and multi-oriented text in natural scene images","date":"2020-11-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-mechanism-transformers-bert-and-gpt","title":"Attention Mechanism, Transformers, BERT, and GPT: Tutorial and Survey","date":"2020-11-17","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-calibration-in-deep-metric-learning","title":"Improving Calibration in Deep Metric Learning With Cross-Example Softmax","date":"2020-11-17","arxiv_id":"2011.08824","n_code_links":0,"syntology":null},{"paper":null,"slug":"mvp-bert-redesigning-vocabularies-for-chinese-1","title":"MVP-BERT: Redesigning Vocabularies for Chinese BERT and Multi-Vocab Pretraining","date":"2020-11-17","arxiv_id":"2011.08539","n_code_links":0,"syntology":null},{"paper":null,"slug":"semi-supervised-learning-of-galaxy-morphology","title":"Semi-supervised Learning of Galaxy Morphology using Equivariant Transformer Variational Autoencoders","date":"2020-11-17","arxiv_id":"2011.08714","n_code_links":0,"syntology":null},{"paper":"/paper/siena-stochastic-multi-expert-neural-patcher","slug":"siena-stochastic-multi-expert-neural-patcher","title":"SHIELD: Defending Textual Neural Networks against Multiple Black-Box Adversarial Attacks with Stochastic Multi-Expert Patcher","date":"2020-11-17","arxiv_id":"2011.08908","n_code_links":1,"syntology":null},{"paper":"/paper/beyond-i-i-d-three-levels-of-generalization","slug":"beyond-i-i-d-three-levels-of-generalization","title":"Beyond I.I.D.: Three Levels of Generalization for Question Answering on Knowledge Bases","date":"2020-11-16","arxiv_id":"2011.07743","n_code_links":1,"syntology":null},{"paper":null,"slug":"don-t-patronize-me-an-annotated-dataset-with","title":"Don't Patronize Me! An Annotated Dataset with Patronizing and Condescending Language towards Vulnerable Communities","date":"2020-11-16","arxiv_id":"2011.08320","n_code_links":0,"syntology":null},{"paper":null,"slug":"frdet-balanced-and-lightweight-object","title":"FRDet: Balanced and Lightweight Object Detector based on Fire-Residual Modules for Embedded Processor of Autonomous Driving","date":"2020-11-16","arxiv_id":"2011.08061","n_code_links":0,"syntology":null},{"paper":null,"slug":"iit-kgp-at-fincausal-2020-shared-task-1","title":"IIT_kgp at FinCausal 2020, Shared Task 1: Causality Detection using Sentence Embeddings in Financial Reports","date":"2020-11-16","arxiv_id":"2011.07670","n_code_links":0,"syntology":null},{"paper":"/paper/it-s-a-thin-line-between-love-and-hate-using","slug":"it-s-a-thin-line-between-love-and-hate-using","title":"It's a Thin Line Between Love and Hate: Using the Echo in Modeling Dynamics of Racist Online Communities","date":"2020-11-16","arxiv_id":"2012.01133","n_code_links":1,"syntology":null},{"paper":"/paper/learning-from-task-descriptions","slug":"learning-from-task-descriptions","title":"Learning from Task Descriptions","date":"2020-11-16","arxiv_id":"2011.08115","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":1,"n_instrument":1,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":null}},{"paper":"/paper/on-the-effectiveness-of-vision-transformers","slug":"on-the-effectiveness-of-vision-transformers","title":"On the Effectiveness of Vision Transformers for Zero-shot Face Anti-Spoofing","date":"2020-11-16","arxiv_id":"2011.08019","n_code_links":1,"syntology":null},{"paper":"/paper/scaled-yolov4-scaling-cross-stage-partial","slug":"scaled-yolov4-scaling-cross-stage-partial","title":"Scaled-YOLOv4: Scaling Cross Stage Partial Network","date":"2020-11-16","arxiv_id":"2011.08036","n_code_links":41,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["WongKinYiu/ScaledYOLOv4"],"state":"official: no sample here; runs from other or unrecorded repositories","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["unlocated"]}}},{"paper":"/paper/real-time-polyp-detection-localisation-and","slug":"real-time-polyp-detection-localisation-and","title":"Real-Time Polyp Detection, Localization and Segmentation in Colonoscopy Using Deep Learning","date":"2020-11-15","arxiv_id":"2011.07631","n_code_links":1,"syntology":null},{"paper":"/paper/actbert-learning-global-local-video-text-1","slug":"actbert-learning-global-local-video-text-1","title":"ActBERT: Learning Global-Local Video-Text Representations","date":"2020-11-14","arxiv_id":"2011.07231","n_code_links":1,"syntology":null},{"paper":"/paper/cl-ims-diacr-ita-volente-o-nolente-bert-does","slug":"cl-ims-diacr-ita-volente-o-nolente-bert-does","title":"CL-IMS @ DIACR-Ita: Volente o Nolente: BERT does not outperform SGNS on Semantic Change Detection","date":"2020-11-14","arxiv_id":"2011.07247","n_code_links":1,"syntology":null},{"paper":"/paper/debatesum-a-large-scale-argument-mining-and","slug":"debatesum-a-large-scale-argument-mining-and","title":"DebateSum: A large-scale argument mining and summarization dataset","date":"2020-11-14","arxiv_id":"2011.07251","n_code_links":3,"syntology":{"ran":1,"of":1,"n_ran_checked":0,"n_instrument":1,"unverified":0,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","official":{"repos":["Hellisotherpeople/DebateSum","Hellisotherpeople/debate2vec","arvind-balaji/debate-cards"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"g-rcn-optimizing-the-gap-between","title":"G-RCN: Optimizing the Gap between Classification and Localization Tasks for Object Detection","date":"2020-11-14","arxiv_id":"2012.03677","n_code_links":0,"syntology":null},{"paper":"/paper/utilizing-bidirectional-encoder","slug":"utilizing-bidirectional-encoder","title":"Utilizing Bidirectional Encoder Representations from Transformers for Answer Selection","date":"2020-11-14","arxiv_id":"2011.07208","n_code_links":1,"syntology":null},{"paper":"/paper/editor-an-edit-based-transformer-with","slug":"editor-an-edit-based-transformer-with","title":"EDITOR: an Edit-Based Transformer with Repositioning for Neural Machine Translation with Soft Lexical Constraints","date":"2020-11-13","arxiv_id":"2011.06868","n_code_links":1,"syntology":null},{"paper":"/paper/flert-document-level-features-for-named","slug":"flert-document-level-features-for-named","title":"FLERT: Document-Level Features for Named Entity Recognition","date":"2020-11-13","arxiv_id":"2011.06993","n_code_links":1,"syntology":null},{"paper":null,"slug":"metastatic-cancer-image-classification-based","title":"Metastatic Cancer Image Classification Based On Deep Learning Method","date":"2020-11-13","arxiv_id":"2011.06984","n_code_links":0,"syntology":null},{"paper":null,"slug":"multi-modal-emotion-detection-with-transfer","title":"Multi-Modal Emotion Detection with Transfer Learning","date":"2020-11-13","arxiv_id":"2011.07065","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-interpretable-end-to-end-fine-tuning","title":"An Interpretable End-to-end Fine-tuning Approach for Long Clinical Text","date":"2020-11-12","arxiv_id":"2011.06504","n_code_links":0,"syntology":null},{"paper":null,"slug":"augmenting-bert-carefully-with","title":"Augmenting BERT Carefully with Underrepresented Linguistic Features","date":"2020-11-12","arxiv_id":"2011.06153","n_code_links":0,"syntology":null},{"paper":"/paper/author-s-sentiment-prediction","slug":"author-s-sentiment-prediction","title":"Author's Sentiment Prediction","date":"2020-11-12","arxiv_id":"2011.06128","n_code_links":1,"syntology":null},{"paper":"/paper/biomedical-named-entity-recognition-at-scale","slug":"biomedical-named-entity-recognition-at-scale","title":"Biomedical Named Entity Recognition at Scale","date":"2020-11-12","arxiv_id":"2011.06315","n_code_links":1,"syntology":null},{"paper":"/paper/dsam-a-distance-shrinking-with-angular","slug":"dsam-a-distance-shrinking-with-angular","title":"DSAM: A Distance Shrinking with Angular Marginalizing Loss for High Performance Vehicle Re-identificatio","date":"2020-11-12","arxiv_id":"2011.06228","n_code_links":0,"syntology":null},{"paper":null,"slug":"identifying-depressive-symptoms-from-tweets","title":"Identifying Depressive Symptoms from Tweets: Figurative Language Enabled Multitask Learning Framework","date":"2020-11-12","arxiv_id":"2011.06149","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-model-accuracy-for-imbalanced-image","title":"Improving Model Accuracy for Imbalanced Image Classification Tasks by Adding a Final Batch Normalization Layer: An Empirical Study","date":"2020-11-12","arxiv_id":"2011.06319","n_code_links":0,"syntology":null},{"paper":null,"slug":"an-ensemble-based-approach-by-fine-tuning-the","title":"An ensemble-based approach by fine-tuning the deep transfer learning models to classify pneumonia from chest X-ray images","date":"2020-11-11","arxiv_id":"2011.05543","n_code_links":0,"syntology":null},{"paper":"/paper/automatic-open-world-reliability-assessment","slug":"automatic-open-world-reliability-assessment","title":"Automatic Open-World Reliability Assessment","date":"2020-11-11","arxiv_id":"2011.05506","n_code_links":1,"syntology":null},{"paper":"/paper/deepi2i-enabling-deep-hierarchical-image-to","slug":"deepi2i-enabling-deep-hierarchical-image-to","title":"DeepI2I: Enabling Deep Hierarchical Image-to-Image Translation by Transferring from GANs","date":"2020-11-11","arxiv_id":"2011.05867","n_code_links":1,"syntology":null},{"paper":"/paper/hurricane-forecasting-a-novel-multimodal","slug":"hurricane-forecasting-a-novel-multimodal","title":"Hurricane Forecasting: A Novel Multimodal Machine Learning Framework","date":"2020-11-11","arxiv_id":"2011.06125","n_code_links":3,"syntology":{"ran":3,"of":4,"n_ran_checked":2,"n_instrument":1,"unverified":1,"pointer_only":4,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["leobix/hurricast"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["listed","official"]}}},{"paper":"/paper/multilingual-irony-detection-with-dependency","slug":"multilingual-irony-detection-with-dependency","title":"Multilingual Irony Detection with Dependency Syntax and Neural Models","date":"2020-11-11","arxiv_id":"2011.05706","n_code_links":1,"syntology":null},{"paper":null,"slug":"nit-covid-19-at-wnut-2020-task-2-deep","title":"NIT COVID-19 at WNUT-2020 Task 2: Deep Learning Model RoBERTa for Identify Informative COVID-19 English Tweets","date":"2020-11-11","arxiv_id":"2011.05551","n_code_links":0,"syntology":null},{"paper":null,"slug":"recognizing-more-emotions-with-less-data","title":"Recognizing More Emotions with Less Data Using Self-supervised Transfer Learning","date":"2020-11-11","arxiv_id":"2011.05585","n_code_links":0,"syntology":null},{"paper":null,"slug":"towards-semi-supervised-semantics","title":"Towards Semi-Supervised Semantics Understanding from Speech","date":"2020-11-11","arxiv_id":"2011.06195","n_code_links":0,"syntology":null},{"paper":null,"slug":"trailer-transformer-based-time-wise-long-term","title":"TERMCast: Temporal Relation Modeling for Effective Urban Flow Forecasting","date":"2020-11-11","arxiv_id":"2011.05554","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-multi-plant-disease-diagnosis-method-using","title":"A Multi-Plant Disease Diagnosis Method using Convolutional Neural Network","date":"2020-11-10","arxiv_id":"2011.05151","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-systematic-comparison-of-encrypted-machine","title":"A Systematic Comparison of Encrypted Machine Learning Solutions for Image Classification","date":"2020-11-10","arxiv_id":"2011.05296","n_code_links":0,"syntology":null},{"paper":null,"slug":"detecting-social-media-manipulation-in-low","title":"Detecting Social Media Manipulation in Low-Resource Languages","date":"2020-11-10","arxiv_id":"2011.05367","n_code_links":0,"syntology":null},{"paper":"/paper/dirichlet-pruning-for-neural-network","slug":"dirichlet-pruning-for-neural-network","title":"Dirichlet Pruning for Neural Network Compression","date":"2020-11-10","arxiv_id":"2011.05985","n_code_links":1,"syntology":null},{"paper":null,"slug":"e-t-entity-transformers-coreference-augmented","title":"E.T.: Entity-Transformers. Coreference augmented Neural Language Model for richer mention representations via Entity-Transformer blocks","date":"2020-11-10","arxiv_id":"2011.05431","n_code_links":0,"syntology":null},{"paper":null,"slug":"ellipse-detection-and-localization-with","title":"Ellipse Detection and Localization with Applications to Knots in Sawn Lumber Images","date":"2020-11-10","arxiv_id":"2011.04844","n_code_links":0,"syntology":null},{"paper":null,"slug":"perturbation-based-exploration-methods-in","title":"Perturbation-based exploration methods in deep reinforcement learning","date":"2020-11-10","arxiv_id":"2011.05446","n_code_links":0,"syntology":null},{"paper":null,"slug":"pretraining-strategies-waveform-model-choice","title":"Pretraining Strategies, Waveform Model Choice, and Acoustic Configurations for Multi-Speaker End-to-End Speech Synthesis","date":"2020-11-10","arxiv_id":"2011.04839","n_code_links":0,"syntology":null},{"paper":null,"slug":"temporal-stochastic-softmax-for-3d-cnns-an","title":"Temporal Stochastic Softmax for 3D CNNs: An Application in Facial Expression Recognition","date":"2020-11-10","arxiv_id":"2011.05227","n_code_links":0,"syntology":null},{"paper":"/paper/umberto-mtsa-accompl-it-improving-complexity","slug":"umberto-mtsa-accompl-it-improving-complexity","title":"UmBERTo-MTSA @ AcCompl-It: Improving Complexity and Acceptability Prediction with Multi-task Learning on Self-Supervised Annotations","date":"2020-11-10","arxiv_id":"2011.05197","n_code_links":1,"syntology":null},{"paper":"/paper/when-do-you-need-billions-of-words-of","slug":"when-do-you-need-billions-of-words-of","title":"When Do You Need Billions of Words of Pretraining Data?","date":"2020-11-10","arxiv_id":"2011.04946","n_code_links":1,"syntology":null},{"paper":null,"slug":"an-improved-helmet-detection-method-for","title":"An improved helmet detection method for YOLOv3 on an unbalanced dataset","date":"2020-11-09","arxiv_id":"2011.04214","n_code_links":0,"syntology":null},{"paper":"/paper/bangla-text-classification-using-transformers","slug":"bangla-text-classification-using-transformers","title":"Bangla Text Classification using Transformers","date":"2020-11-09","arxiv_id":"2011.04446","n_code_links":1,"syntology":null},{"paper":null,"slug":"bert-jam-boosting-bert-enhanced-neural","title":"BERT-JAM: Boosting BERT-Enhanced Neural Machine Translation with Joint Attention","date":"2020-11-09","arxiv_id":"2011.04266","n_code_links":0,"syntology":null},{"paper":null,"slug":"catch-the-tails-of-bert","title":"Positional Artefacts Propagate Through Masked Language Model Embeddings","date":"2020-11-09","arxiv_id":"2011.04393","n_code_links":0,"syntology":null},{"paper":"/paper/cxgbert-bert-meets-construction-grammar","slug":"cxgbert-bert-meets-construction-grammar","title":"CxGBERT: BERT meets Construction Grammar","date":"2020-11-09","arxiv_id":"2011.04134","n_code_links":1,"syntology":null},{"paper":null,"slug":"estbert-a-pretrained-language-specific-bert","title":"EstBERT: A Pretrained Language-Specific BERT for Estonian","date":"2020-11-09","arxiv_id":"2011.04784","n_code_links":0,"syntology":null},{"paper":"/paper/language-through-a-prism-a-spectral-approach","slug":"language-through-a-prism-a-spectral-approach","title":"Language Through a Prism: A Spectral Approach for Multiscale Language Representations","date":"2020-11-09","arxiv_id":"2011.04823","n_code_links":1,"syntology":{"ran":0,"of":5,"n_ran_checked":0,"n_instrument":0,"unverified":5,"pointer_only":0,"phrase":"0 ran · 5 unverified","official":null}},{"paper":"/paper/magneto-an-efficient-deep-learning-method-for-1","slug":"magneto-an-efficient-deep-learning-method-for-1","title":"MAGNeto: An Efficient Deep Learning Method for the Extractive Tags Summarization Problem","date":"2020-11-09","arxiv_id":"2011.04349","n_code_links":1,"syntology":null},{"paper":"/paper/real-time-object-detection-method-based-on","slug":"real-time-object-detection-method-based-on","title":"Real-time object detection method based on improved YOLOv4-tiny","date":"2020-11-09","arxiv_id":"2011.04244","n_code_links":1,"syntology":null},{"paper":null,"slug":"superdeconfuse-a-supervised-deep","title":"SuperDeConFuse: A Supervised Deep Convolutional Transform based Fusion Framework for Financial Trading Systems","date":"2020-11-09","arxiv_id":"2011.04364","n_code_links":0,"syntology":null},{"paper":"/paper/visbert-hidden-state-visualizations-for","slug":"visbert-hidden-state-visualizations-for","title":"VisBERT: Hidden-State Visualizations for Transformers","date":"2020-11-09","arxiv_id":"2011.04507","n_code_links":1,"syntology":null},{"paper":"/paper/adapting-a-language-model-for-controlled","slug":"adapting-a-language-model-for-controlled","title":"Adapting a Language Model for Controlled Affective Text Generation","date":"2020-11-08","arxiv_id":"2011.04000","n_code_links":1,"syntology":{"ran":2,"of":3,"n_ran_checked":1,"n_instrument":1,"unverified":1,"pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 1 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 1 unverified","official":{"repos":["ishikasingh/Affective-text-gen"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official","unlocated"]}}},{"paper":"/paper/indicnlpsuite-monolingual-corpora-evaluation","slug":"indicnlpsuite-monolingual-corpora-evaluation","title":"IndicNLPSuite: Monolingual Corpora, Evaluation Benchmarks and Pre-trained Multilingual Language Models for Indian Languages","date":"2020-11-08","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/long-range-arena-a-benchmark-for-efficient-1","slug":"long-range-arena-a-benchmark-for-efficient-1","title":"Long Range Arena: A Benchmark for Efficient Transformers","date":"2020-11-08","arxiv_id":"2011.04006","n_code_links":5,"syntology":{"ran":2,"of":3,"n_ran_checked":2,"n_instrument":0,"unverified":1,"pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","official":{"repos":["google-research/long-range-arena"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"predictive-analysis-of-diabetic-retinopathy","title":"Predictive Analysis of Diabetic Retinopathy with Transfer Learning","date":"2020-11-08","arxiv_id":"2011.04052","n_code_links":0,"syntology":null},{"paper":"/paper/stochastic-attention-head-removal-a-simple","slug":"stochastic-attention-head-removal-a-simple","title":"Stochastic Attention Head Removal: A simple and effective method for improving Transformer Based ASR Models","date":"2020-11-08","arxiv_id":"2011.04004","n_code_links":5,"syntology":{"ran":0,"of":1,"n_ran_checked":0,"n_instrument":0,"unverified":1,"pointer_only":1,"phrase":"0 ran · 1 unverified","official":{"repos":["s1603602/attention_head_removal"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}},{"paper":"/paper/a-reinforcement-learning-approach-to-the","slug":"a-reinforcement-learning-approach-to-the","title":"A Reinforcement Learning Approach to the Orienteering Problem with Time Windows","date":"2020-11-07","arxiv_id":"2011.03647","n_code_links":2,"syntology":null},{"paper":null,"slug":"know-what-you-don-t-need-single-shot-meta","title":"Know What You Don't Need: Single-Shot Meta-Pruning for Attention Heads","date":"2020-11-07","arxiv_id":"2011.03770","n_code_links":0,"syntology":null},{"paper":null,"slug":"rethinking-the-value-of-transformer","title":"Rethinking the Value of Transformer Components","date":"2020-11-07","arxiv_id":"2011.03803","n_code_links":0,"syntology":null},{"paper":"/paper/seqgensql-a-robust-sequence-generation-model","slug":"seqgensql-a-robust-sequence-generation-model","title":"SeqGenSQL -- A Robust Sequence Generation Model for Structured Query Language","date":"2020-11-07","arxiv_id":"2011.03836","n_code_links":2,"syntology":null},{"paper":null,"slug":"channel-pruning-via-multi-criteria-based-on","title":"Channel Pruning via Multi-Criteria based on Weight Dependency","date":"2020-11-06","arxiv_id":"2011.03240","n_code_links":0,"syntology":null},{"paper":"/paper/deep-transfer-learning-for-automated","slug":"deep-transfer-learning-for-automated","title":"Deep Transfer Learning for Automated Diagnosis of Skin Lesions from Photographs","date":"2020-11-06","arxiv_id":"2011.04475","n_code_links":1,"syntology":null},{"paper":"/paper/from-dataset-recycling-to-multi-property","slug":"from-dataset-recycling-to-multi-property","title":"From Dataset Recycling to Multi-Property Extraction and Beyond","date":"2020-11-06","arxiv_id":"2011.03228","n_code_links":1,"syntology":null},{"paper":null,"slug":"highly-available-data-parallel-ml-training-on","title":"Highly Available Data Parallel ML training on Mesh Networks","date":"2020-11-06","arxiv_id":"2011.03605","n_code_links":0,"syntology":null},{"paper":null,"slug":"improving-prosody-modelling-with-cross","title":"Improving Prosody Modelling with Cross-Utterance BERT Embeddings for End-to-end Speech Synthesis","date":"2020-11-06","arxiv_id":"2011.05161","n_code_links":0,"syntology":null},{"paper":"/paper/semi-supervised-low-resource-style-transfer","slug":"semi-supervised-low-resource-style-transfer","title":"Semi-Supervised Low-Resource Style Transfer of Indonesian Informal to Formal Language with Iterative Forward-Translation","date":"2020-11-06","arxiv_id":"2011.03286","n_code_links":1,"syntology":null},{"paper":null,"slug":"bw-eda-eend-streaming-end-to-end-neural","title":"BW-EDA-EEND: Streaming End-to-End Neural Speaker Diarization for a Variable Number of Speakers","date":"2020-11-05","arxiv_id":"2011.02678","n_code_links":0,"syntology":null}],"record_sha256":"69209545c105db9fcc80c92a7c67460975edcb76281fc04eb0eff3a449b9347e","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}