{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/language-modeling/papers/56","list_of":"/task/language-modeling","task":"Language Modeling","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":56,"pages_in_order":142,"rows_per_page":100,"rows":[5501,5600],"of":14182,"counts":{"archive_papers_tagged":14182,"with_a_code_link":5620,"where_syntology_ran_a_sample":1894,"not_listed_spam_title":0,"listed":14182,"listed_where_code_ran":1894,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":1580,"every_run_a_failure_of_syntologys_instrument":314,"listed_with_a_run_with_no_instrument_failure":1580,"listed_every_run_a_failure_of_syntologys_instrument":314,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/language-modeling","prev":"/task/language-modeling/papers/55","next":"/task/language-modeling/papers/57","papers":[{"url":"/paper/a-context-based-approach-for-dialogue-act","slug":"a-context-based-approach-for-dialogue-act","title":"A Context-based Approach for Dialogue Act Recognition using Simple Recurrent Neural Networks","date":"2018-05-16","arxiv_id":"1805.06280","repositories_listed":1,"syntology":null},{"url":"/paper/polite-dialogue-generation-without-parallel","slug":"polite-dialogue-generation-without-parallel","title":"Polite Dialogue Generation Without Parallel Data","date":"2018-05-08","arxiv_id":"1805.03162","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":1,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/polite-dialogue-generation-without-parallel#ran","syntology_url":"https://syntology.ai/paper/1805.03162","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1805.03162"}},"official":null}},{"url":"/paper/hierarchical-pointer-generator-memory-network","slug":"hierarchical-pointer-generator-memory-network","title":"Disentangling Language and Knowledge in Task-Oriented Dialogs","date":"2018-05-03","arxiv_id":"1805.01216","repositories_listed":1,"syntology":null},{"url":"/paper/diacritics-restoration-using-neural-networks","slug":"diacritics-restoration-using-neural-networks","title":"Diacritics Restoration Using Neural Networks","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/evaluation-phonemic-transcription-of-low","slug":"evaluation-phonemic-transcription-of-low","title":"Evaluation Phonemic Transcription of Low-Resource Tonal Languages for Language Documentation","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/fonbund-a-library-for-combining-cross-lingual","slug":"fonbund-a-library-for-combining-cross-lingual","title":"FonBund: A Library for Combining Cross-lingual Phonological Segment Data","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/preparation-and-usage-of-xhosa","slug":"preparation-and-usage-of-xhosa","title":"Preparation and Usage of Xhosa Lexicographical Data for a Multilingual, Federated Environment","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/tf-lm-tensorflow-based-language-modeling","slug":"tf-lm-tensorflow-based-language-modeling","title":"TF-LM: TensorFlow-based Language Modeling Toolkit","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/syllable-based-sequence-to-sequence-speech","slug":"syllable-based-sequence-to-sequence-speech","title":"Syllable-Based Sequence-to-Sequence Speech Recognition with the Transformer in Mandarin Chinese","date":"2018-04-28","arxiv_id":"1804.10752","repositories_listed":1,"syntology":null},{"url":"/paper/spell-once-summon-anywhere-a-two-level-open","slug":"spell-once-summon-anywhere-a-two-level-open","title":"Spell Once, Summon Anywhere: A Two-Level Open-Vocabulary Language Model","date":"2018-04-23","arxiv_id":"1804.08205","repositories_listed":1,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/spell-once-summon-anywhere-a-two-level-open#ran","syntology_url":"https://syntology.ai/paper/1804.08205","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1804.08205"}},"official":{"repos":["sjmielke/spell-once"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/efficient-contextualized-representation","slug":"efficient-contextualized-representation","title":"Efficient Contextualized Representation: Language Model Pruning for Sequence Labeling","date":"2018-04-20","arxiv_id":"1804.07827","repositories_listed":1,"syntology":null},{"url":"/paper/identifying-compromised-accounts-on-social","slug":"identifying-compromised-accounts-on-social","title":"Semantic Text Analysis for Detection of Compromised Accounts on Social Networks","date":"2018-04-19","arxiv_id":"1804.07247","repositories_listed":1,"syntology":null},{"url":"/paper/automated-evaluation-of-out-of-context-errors","slug":"automated-evaluation-of-out-of-context-errors","title":"Automated Evaluation of Out-of-Context Errors","date":"2018-03-23","arxiv_id":"1803.08983","repositories_listed":1,"syntology":null},{"url":"/paper/neural-lattice-language-models","slug":"neural-lattice-language-models","title":"Neural Lattice Language Models","date":"2018-03-13","arxiv_id":"1803.05071","repositories_listed":1,"syntology":null},{"url":"/paper/the-importance-of-being-recurrent-for","slug":"the-importance-of-being-recurrent-for","title":"The Importance of Being Recurrent for Modeling Hierarchical Structure","date":"2018-03-09","arxiv_id":"1803.03585","repositories_listed":1,"syntology":null},{"url":"/paper/deep-fsmn-for-large-vocabulary-continuous","slug":"deep-fsmn-for-large-vocabulary-continuous","title":"Deep-FSMN for Large Vocabulary Continuous Speech Recognition","date":"2018-03-04","arxiv_id":"1803.05030","repositories_listed":1,"syntology":null},{"url":"/paper/recurrent-neural-network-based-semantic","slug":"recurrent-neural-network-based-semantic","title":"Recurrent Neural Network-Based Semantic Variational Autoencoder for Sequence-to-Sequence Learning","date":"2018-02-09","arxiv_id":"1802.03238","repositories_listed":1,"syntology":null},{"url":"/paper/learning-from-past-mistakes-improving","slug":"learning-from-past-mistakes-improving","title":"Learning from Past Mistakes: Improving Automatic Speech Recognition Output via Noisy-Clean Phrase Context Modeling","date":"2018-02-07","arxiv_id":"1802.02607","repositories_listed":1,"syntology":null},{"url":"/paper/nested-lstms","slug":"nested-lstms","title":"Nested LSTMs","date":"2018-01-31","arxiv_id":"1801.10308","repositories_listed":1,"syntology":null},{"url":"/paper/strassennets-deep-learning-with-a","slug":"strassennets-deep-learning-with-a","title":"StrassenNets: Deep Learning with a Multiplication Budget","date":"2017-12-11","arxiv_id":"1712.03942","repositories_listed":1,"syntology":null},{"url":"/paper/contextualized-word-representations-for","slug":"contextualized-word-representations-for","title":"Contextualized Word Representations for Reading Comprehension","date":"2017-12-10","arxiv_id":"1712.03609","repositories_listed":1,"syntology":null},{"url":"/paper/phonemic-transcription-of-low-resource-tonal","slug":"phonemic-transcription-of-low-resource-tonal","title":"Phonemic Transcription of Low-Resource Tonal Languages","date":"2017-12-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/text-generation-based-on-generative","slug":"text-generation-based-on-generative","title":"Text Generation Based on Generative Adversarial Nets with Latent Variable","date":"2017-12-01","arxiv_id":"1712.00170","repositories_listed":1,"syntology":null},{"url":"/paper/z-forcing-training-stochastic-recurrent","slug":"z-forcing-training-stochastic-recurrent","title":"Z-Forcing: Training Stochastic Recurrent Networks","date":"2017-11-15","arxiv_id":"1711.05411","repositories_listed":1,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":3,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":5,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/z-forcing-training-stochastic-recurrent#ran","syntology_url":"https://syntology.ai/paper/1711.05411","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1711.05411"}},"official":{"repos":["anirudh9119/zforcing_nips17"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/rdf2vec-rdf-graph-embeddings-and-their","slug":"rdf2vec-rdf-graph-embeddings-and-their","title":"RDF2Vec: RDF Graph Embeddings and Their Applications","date":"2017-11-10","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/neural-language-modeling-by-jointly-learning","slug":"neural-language-modeling-by-jointly-learning","title":"Neural Language Modeling by Jointly Learning Syntax and Lexicon","date":"2017-11-02","arxiv_id":"1711.02013","repositories_listed":1,"syntology":{"n":5,"n_ran":4,"n_constructed":0,"n_ran_checked":4,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/neural-language-modeling-by-jointly-learning#ran","syntology_url":"https://syntology.ai/paper/1711.02013","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1711.02013"}},"official":null}},{"url":"/paper/improving-low-resource-neural-machine","slug":"improving-low-resource-neural-machine","title":"Improving Low-Resource Neural Machine Translation with Filtered Pseudo-Parallel Corpus","date":"2017-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/kyoto-university-participation-to-wat-2017","slug":"kyoto-university-participation-to-wat-2017","title":"Kyoto University Participation to WAT 2017","date":"2017-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/fraternal-dropout","slug":"fraternal-dropout","title":"Fraternal Dropout","date":"2017-10-31","arxiv_id":"1711.00066","repositories_listed":1,"syntology":null},{"url":"/paper/a-dual-encoder-sequence-to-sequence-model-for","slug":"a-dual-encoder-sequence-to-sequence-model-for","title":"A Dual Encoder Sequence to Sequence Model for Open-Domain Dialogue Modeling","date":"2017-10-28","arxiv_id":"1710.10520","repositories_listed":1,"syntology":null},{"url":"/paper/the-implementation-of-a-deep-recurrent-neural","slug":"the-implementation-of-a-deep-recurrent-neural","title":"The implementation of a Deep Recurrent Neural Network Language Model on a Xilinx FPGA","date":"2017-10-26","arxiv_id":"1710.10296","repositories_listed":1,"syntology":null},{"url":"/paper/low-rank-rnn-adaptation-for-context-aware","slug":"low-rank-rnn-adaptation-for-context-aware","title":"Low-Rank RNN Adaptation for Context-Aware Language Modeling","date":"2017-10-06","arxiv_id":"1710.02603","repositories_listed":1,"syntology":null},{"url":"/paper/counterfactual-language-model-adaptation-for","slug":"counterfactual-language-model-adaptation-for","title":"Counterfactual Language Model Adaptation for Suggesting Phrases","date":"2017-10-04","arxiv_id":"1710.01799","repositories_listed":1,"syntology":null},{"url":"/paper/shifting-mean-activation-towards-zero-with","slug":"shifting-mean-activation-towards-zero-with","title":"Shifting Mean Activation Towards Zero with Bipolar Activation Functions","date":"2017-09-12","arxiv_id":"1709.04054","repositories_listed":1,"syntology":null},{"url":"/paper/cynical-selection-of-language-model-training","slug":"cynical-selection-of-language-model-training","title":"Cynical Selection of Language Model Training Data","date":"2017-09-07","arxiv_id":"1709.02279","repositories_listed":1,"syntology":null},{"url":"/paper/a-neural-language-model-for-dynamically","slug":"a-neural-language-model-for-dynamically","title":"A Neural Language Model for Dynamically Representing the Meanings of Unknown Words and Entities in a Discourse","date":"2017-09-06","arxiv_id":"1709.01679","repositories_listed":1,"syntology":null},{"url":"/paper/patterns-versus-characters-in-subword-aware","slug":"patterns-versus-characters-in-subword-aware","title":"Patterns versus Characters in Subword-aware Neural Language Modeling","date":"2017-09-02","arxiv_id":"1709.00541","repositories_listed":1,"syntology":null},{"url":"/paper/using-target-side-monolingual-data-for-neural","slug":"using-target-side-monolingual-data-for-neural","title":"Using Target-side Monolingual Data for Neural Machine Translation through Multi-task Learning","date":"2017-09-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/glyph-aware-embedding-of-chinese-characters","slug":"glyph-aware-embedding-of-chinese-characters","title":"Glyph-aware Embedding of Chinese Characters","date":"2017-08-31","arxiv_id":"1709.00028","repositories_listed":1,"syntology":null},{"url":"/paper/gradual-learning-of-recurrent-neural-networks","slug":"gradual-learning-of-recurrent-neural-networks","title":"Gradual Learning of Recurrent Neural Networks","date":"2017-08-29","arxiv_id":"1708.08863","repositories_listed":1,"syntology":null},{"url":"/paper/vqs-linking-segmentations-to-questions-and","slug":"vqs-linking-segmentations-to-questions-and","title":"VQS: Linking Segmentations to Questions and Answers for Supervised Attention in VQA and Question-Focused Semantic Segmentation","date":"2017-08-15","arxiv_id":"1708.04686","repositories_listed":1,"syntology":null},{"url":"/paper/location-name-extraction-from-targeted-text","slug":"location-name-extraction-from-targeted-text","title":"Location Name Extraction from Targeted Text Streams using Gazetteer-based Statistical Language Models","date":"2017-08-10","arxiv_id":"1708.03105","repositories_listed":1,"syntology":null},{"url":"/paper/detecting-anxiety-through-reddit","slug":"detecting-anxiety-through-reddit","title":"Detecting Anxiety through Reddit","date":"2017-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/transition-based-generation-from-abstract","slug":"transition-based-generation-from-abstract","title":"Transition-Based Generation from Abstract Meaning Representations","date":"2017-07-24","arxiv_id":"1707.07591","repositories_listed":1,"syntology":null},{"url":"/paper/high-risk-learning-acquiring-new-word-vectors","slug":"high-risk-learning-acquiring-new-word-vectors","title":"High-risk learning: acquiring new word vectors from tiny data","date":"2017-07-20","arxiv_id":"1707.06556","repositories_listed":1,"syntology":null},{"url":"/paper/syllable-aware-neural-language-models-a","slug":"syllable-aware-neural-language-models-a","title":"Syllable-aware Neural Language Models: A Failure to Beat Character-aware Ones","date":"2017-07-20","arxiv_id":"1707.06480","repositories_listed":1,"syntology":null},{"url":"/paper/an-embedded-deep-learning-based-word","slug":"an-embedded-deep-learning-based-word","title":"An Embedded Deep Learning based Word Prediction","date":"2017-07-06","arxiv_id":"1707.01662","repositories_listed":1,"syntology":null},{"url":"/paper/improved-word-representation-learning-with","slug":"improved-word-representation-learning-with","title":"Improved Word Representation Learning with Sememes","date":"2017-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/joint-ctcattention-decoding-for-end-to-end","slug":"joint-ctcattention-decoding-for-end-to-end","title":"Joint CTC/attention decoding for end-to-end speech recognition","date":"2017-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/device-placement-optimization-with","slug":"device-placement-optimization-with","title":"Device Placement Optimization with Reinforcement Learning","date":"2017-06-13","arxiv_id":"1706.04972","repositories_listed":1,"syntology":null},{"url":"/paper/exploring-the-syntactic-abilities-of-rnns","slug":"exploring-the-syntactic-abilities-of-rnns","title":"Exploring the Syntactic Abilities of RNNs with Multi-task Learning","date":"2017-06-12","arxiv_id":"1706.03542","repositories_listed":1,"syntology":null},{"url":"/paper/biased-importance-sampling-for-deep-neural","slug":"biased-importance-sampling-for-deep-neural","title":"Biased Importance Sampling for Deep Neural Network Training","date":"2017-05-31","arxiv_id":"1706.00043","repositories_listed":1,"syntology":null},{"url":"/paper/fast-slow-recurrent-neural-networks","slug":"fast-slow-recurrent-neural-networks","title":"Fast-Slow Recurrent Neural Networks","date":"2017-05-24","arxiv_id":"1705.08639","repositories_listed":1,"syntology":null},{"url":"/paper/generating-memorable-mnemonic-encodings-of","slug":"generating-memorable-mnemonic-encodings-of","title":"Generating Memorable Mnemonic Encodings of Numbers","date":"2017-05-07","arxiv_id":"1705.02700","repositories_listed":1,"syntology":null},{"url":"/paper/finnish-resources-for-evaluating-language","slug":"finnish-resources-for-evaluating-language","title":"Finnish resources for evaluating language model semantics","date":"2017-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/topically-driven-neural-language-model","slug":"topically-driven-neural-language-model","title":"Topically Driven Neural Language Model","date":"2017-04-26","arxiv_id":"1704.08012","repositories_listed":1,"syntology":null},{"url":"/paper/improving-context-aware-language-models","slug":"improving-context-aware-language-models","title":"Improving Context Aware Language Models","date":"2017-04-21","arxiv_id":"1704.06380","repositories_listed":1,"syntology":null},{"url":"/paper/social-bias-in-elicited-natural-language","slug":"social-bias-in-elicited-natural-language","title":"Social Bias in Elicited Natural Language Inferences","date":"2017-04-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/data-noising-as-smoothing-in-neural-network","slug":"data-noising-as-smoothing-in-neural-network","title":"Data Noising as Smoothing in Neural Network Language Models","date":"2017-03-07","arxiv_id":"1703.02573","repositories_listed":1,"syntology":null},{"url":"/paper/leveraging-large-amounts-of-weakly-supervised","slug":"leveraging-large-amounts-of-weakly-supervised","title":"Leveraging Large Amounts of Weakly Supervised Data for Multi-Language Sentiment Classification","date":"2017-03-07","arxiv_id":"1703.02504","repositories_listed":1,"syntology":null},{"url":"/paper/dynamic-word-embeddings","slug":"dynamic-word-embeddings","title":"Dynamic Word Embeddings","date":"2017-02-27","arxiv_id":"1702.08359","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/dynamic-word-embeddings#ran","syntology_url":"https://syntology.ai/paper/1702.08359","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1702.08359"}},"official":null}},{"url":"/paper/the-effect-of-different-writing-tasks-on","slug":"the-effect-of-different-writing-tasks-on","title":"The Effect of Different Writing Tasks on Linguistic Style: A Case Study of the ROC Story Cloze Task","date":"2017-02-07","arxiv_id":"1702.01841","repositories_listed":1,"syntology":null},{"url":"/paper/first-automatic-fongbe-continuous-speech","slug":"first-automatic-fongbe-continuous-speech","title":"First Automatic Fongbe Continuous Speech Recognition System: Development of Acoustic Models and Language Models","date":"2017-01-21","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/context-aware-captions-from-context-agnostic","slug":"context-aware-captions-from-context-agnostic","title":"Context-aware Captions from Context-agnostic Supervision","date":"2017-01-11","arxiv_id":"1701.02870","repositories_listed":1,"syntology":null},{"url":"/paper/a-character-word-compositional-neural","slug":"a-character-word-compositional-neural","title":"A Character-Word Compositional Neural Language Model for Finnish","date":"2016-12-10","arxiv_id":"1612.03266","repositories_listed":1,"syntology":null},{"url":"/paper/kyoto-nmt-a-neural-machine-translation","slug":"kyoto-nmt-a-neural-machine-translation","title":"Kyoto-NMT: a Neural Machine Translation implementation in Chainer","date":"2016-12-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/syntactic-realization-with-data-driven-neural","slug":"syntactic-realization-with-data-driven-neural","title":"Syntactic realization with data-driven neural tree grammars","date":"2016-12-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/what-do-recurrent-neural-network-grammars","slug":"what-do-recurrent-neural-network-grammars","title":"What Do Recurrent Neural Network Grammars Learn About Syntax?","date":"2016-11-17","arxiv_id":"1611.05774","repositories_listed":1,"syntology":null},{"url":"/paper/topicrnn-a-recurrent-neural-network-with-long","slug":"topicrnn-a-recurrent-neural-network-with-long","title":"TopicRNN: A Recurrent Neural Network with Long-Range Semantic Dependency","date":"2016-11-05","arxiv_id":"1611.01702","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/topicrnn-a-recurrent-neural-network-with-long#ran","syntology_url":"https://syntology.ai/paper/1611.01702","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1611.01702"}},"official":null}},{"url":"/paper/convolutional-neural-network-language-models","slug":"convolutional-neural-network-language-models","title":"Convolutional Neural Network Language Models","date":"2016-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/dual-learning-for-machine-translation","slug":"dual-learning-for-machine-translation","title":"Dual Learning for Machine Translation","date":"2016-11-01","arxiv_id":"1611.00179","repositories_listed":1,"syntology":null},{"url":"/paper/latent-tree-language-model-1","slug":"latent-tree-language-model-1","title":"Latent Tree Language Model","date":"2016-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/parsing-as-language-modeling","slug":"parsing-as-language-modeling","title":"Parsing as Language Modeling","date":"2016-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/professor-forcing-a-new-algorithm-for","slug":"professor-forcing-a-new-algorithm-for","title":"Professor Forcing: A New Algorithm for Training Recurrent Networks","date":"2016-10-27","arxiv_id":"1610.09038","repositories_listed":1,"syntology":null},{"url":"/paper/compressing-neural-language-models-by-sparse","slug":"compressing-neural-language-models-by-sparse","title":"Compressing Neural Language Models by Sparse Word Representations","date":"2016-10-13","arxiv_id":"1610.03950","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/compressing-neural-language-models-by-sparse#ran","syntology_url":"https://syntology.ai/paper/1610.03950","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1610.03950"}},"official":{"repos":["chenych11/lm"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/language-models-with-pre-trained-glove-word","slug":"language-models-with-pre-trained-glove-word","title":"Language Models with Pre-Trained (GloVe) Word Embeddings","date":"2016-10-12","arxiv_id":"1610.03759","repositories_listed":1,"syntology":null},{"url":"/paper/fast-small-and-exact-infinite-order-language","slug":"fast-small-and-exact-infinite-order-language","title":"Fast, Small and Exact: Infinite-order Language Modelling with Compressed Suffix Trees","date":"2016-08-16","arxiv_id":"1608.04465","repositories_listed":1,"syntology":null},{"url":"/paper/a-deep-language-model-for-software-code","slug":"a-deep-language-model-for-software-code","title":"A deep language model for software code","date":"2016-08-09","arxiv_id":"1608.02715","repositories_listed":1,"syntology":null},{"url":"/paper/latent-tree-language-model","slug":"latent-tree-language-model","title":"Latent Tree Language Model","date":"2016-07-24","arxiv_id":"1607.07057","repositories_listed":1,"syntology":null},{"url":"/paper/watch-what-you-just-said-image-captioning","slug":"watch-what-you-just-said-image-captioning","title":"Watch What You Just Said: Image Captioning with Text-Conditional Attention","date":"2016-06-15","arxiv_id":"1606.04621","repositories_listed":1,"syntology":null},{"url":"/paper/learning-to-generate-compositional-color","slug":"learning-to-generate-compositional-color","title":"Learning to Generate Compositional Color Descriptions","date":"2016-06-13","arxiv_id":"1606.03821","repositories_listed":1,"syntology":null},{"url":"/paper/generalizing-and-hybridizing-count-based-and","slug":"generalizing-and-hybridizing-count-based-and","title":"Generalizing and Hybridizing Count-based and Neural Language Models","date":"2016-06-01","arxiv_id":"1606.00499","repositories_listed":1,"syntology":null},{"url":"/paper/temporal-action-detection-using-a-statistical","slug":"temporal-action-detection-using-a-statistical","title":"Temporal Action Detection Using a Statistical Language Model","date":"2016-06-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/transition-based-syntactic-linearization-with","slug":"transition-based-syntactic-linearization-with","title":"Transition-Based Syntactic Linearization with Lookahead Features","date":"2016-06-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/unimelb-at-semeval-2016-tasks-4a-and-4b-an","slug":"unimelb-at-semeval-2016-tasks-4a-and-4b-an","title":"UNIMELB at SemEval-2016 Tasks 4A and 4B: An Ensemble of Neural Networks and a Word2Vec Based Model for Sentiment Classification","date":"2016-06-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/scale-a-scalable-language-engineering-toolkit","slug":"scale-a-scalable-language-engineering-toolkit","title":"SCALE: A Scalable Language Engineering Toolkit","date":"2016-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/the-z-loss-a-shift-and-scale-invariant","slug":"the-z-loss-a-shift-and-scale-invariant","title":"The Z-loss: a shift and scale invariant classification loss belonging to the Spherical Family","date":"2016-04-29","arxiv_id":"1604.08859","repositories_listed":1,"syntology":null},{"url":"/paper/word-ordering-without-syntax","slug":"word-ordering-without-syntax","title":"Word Ordering Without Syntax","date":"2016-04-28","arxiv_id":"1604.08633","repositories_listed":1,"syntology":null},{"url":"/paper/lstm-based-conversation-models","slug":"lstm-based-conversation-models","title":"LSTM based Conversation Models","date":"2016-03-31","arxiv_id":"1603.09457","repositories_listed":1,"syntology":null},{"url":"/paper/a-latent-variable-recurrent-neural-network","slug":"a-latent-variable-recurrent-neural-network","title":"A Latent Variable Recurrent Neural Network for Discourse Relation Language Models","date":"2016-03-07","arxiv_id":"1603.01913","repositories_listed":1,"syntology":null},{"url":"/paper/representation-of-linguistic-form-and","slug":"representation-of-linguistic-form-and","title":"Representation of linguistic form and function in recurrent neural networks","date":"2016-02-29","arxiv_id":"1602.08952","repositories_listed":1,"syntology":null},{"url":"/paper/domain-specific-author-attribution-based-on","slug":"domain-specific-author-attribution-based-on","title":"Domain Specific Author Attribution Based on Feedforward Neural Network Language Models","date":"2016-02-24","arxiv_id":"1602.07393","repositories_listed":1,"syntology":null},{"url":"/paper/on-training-bi-directional-neural-network","slug":"on-training-bi-directional-neural-network","title":"On Training Bi-directional Neural Network Language Model with Noise Contrastive Estimation","date":"2016-02-19","arxiv_id":"1602.06064","repositories_listed":1,"syntology":null},{"url":"/paper/authorship-attribution-using-a-neural-network","slug":"authorship-attribution-using-a-neural-network","title":"Authorship Attribution Using a Neural Network Language Model","date":"2016-02-17","arxiv_id":"1602.05292","repositories_listed":1,"syntology":null},{"url":"/paper/character-level-incremental-speech","slug":"character-level-incremental-speech","title":"Character-Level Incremental Speech Recognition with Recurrent Neural Networks","date":"2016-01-25","arxiv_id":"1601.06581","repositories_listed":1,"syntology":null},{"url":"/paper/regularizing-rnns-by-stabilizing-activations","slug":"regularizing-rnns-by-stabilizing-activations","title":"Regularizing RNNs by Stabilizing Activations","date":"2015-11-26","arxiv_id":"1511.08400","repositories_listed":1,"syntology":null},{"url":"/paper/densecap-fully-convolutional-localization","slug":"densecap-fully-convolutional-localization","title":"DenseCap: Fully Convolutional Localization Networks for Dense Captioning","date":"2015-11-24","arxiv_id":"1511.07571","repositories_listed":1,"syntology":null},{"url":"/paper/blackout-speeding-up-recurrent-neural-network","slug":"blackout-speeding-up-recurrent-neural-network","title":"BlackOut: Speeding up Recurrent Neural Network Language Models With Very Large Vocabularies","date":"2015-11-21","arxiv_id":"1511.06909","repositories_listed":1,"syntology":null},{"url":"/paper/alternative-structures-for-character-level","slug":"alternative-structures-for-character-level","title":"Alternative structures for character-level RNNs","date":"2015-11-19","arxiv_id":"1511.06303","repositories_listed":1,"syntology":null},{"url":"/paper/task-loss-estimation-for-sequence-prediction","slug":"task-loss-estimation-for-sequence-prediction","title":"Task Loss Estimation for Sequence Prediction","date":"2015-11-19","arxiv_id":"1511.06456","repositories_listed":1,"syntology":{"n":1,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"0 ran · 1 unverified","sample_list":"/paper/task-loss-estimation-for-sequence-prediction#ran","syntology_url":"https://syntology.ai/paper/1511.06456","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1511.06456"}},"official":{"repos":["rizar/attention-lvcsr"],"state":"official: harvested, nothing ran","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":[]}}}],"record_sha256":"8dbb424824604241f62a06e5b1cb06232a29ad76d2c6c10591b67b0900337bad","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}