{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/language-modelling/papers/69","list_of":"/task/language-modelling","task":"Language Modelling","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":69,"pages_in_order":177,"rows_per_page":100,"rows":[6801,6900],"of":17610,"counts":{"archive_papers_tagged":17610,"with_a_code_link":7012,"where_syntology_ran_a_sample":2428,"not_listed_spam_title":0,"listed":17610,"listed_where_code_ran":2428,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":2027,"every_run_a_failure_of_syntologys_instrument":401,"listed_with_a_run_with_no_instrument_failure":2027,"listed_every_run_a_failure_of_syntologys_instrument":401,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/language-modelling","prev":"/task/language-modelling/papers/68","next":"/task/language-modelling/papers/70","papers":[{"url":"/paper/zoho-at-semeval-2019-task-9-semi-supervised","slug":"zoho-at-semeval-2019-task-9-semi-supervised","title":"Zoho at SemEval-2019 Task 9: Semi-supervised Domain Adaptation using Tri-training for Suggestion Mining","date":"2019-02-27","arxiv_id":"1902.10623","repositories_listed":1,"syntology":null},{"url":"/paper/polyglot-contextual-representations-improve","slug":"polyglot-contextual-representations-improve","title":"Polyglot Contextual Representations Improve Crosslingual Transfer","date":"2019-02-26","arxiv_id":"1902.09697","repositories_listed":1,"syntology":null},{"url":"/paper/a-fully-differentiable-beam-search-decoder","slug":"a-fully-differentiable-beam-search-decoder","title":"A Fully Differentiable Beam Search Decoder","date":"2019-02-16","arxiv_id":"1902.06022","repositories_listed":1,"syntology":{"n":3,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"0 ran · 3 unverified","sample_list":"/paper/a-fully-differentiable-beam-search-decoder#ran","syntology_url":"https://syntology.ai/paper/1902.06022","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1902.06022"}},"official":null}},{"url":"/paper/the-referential-reader-a-recurrent-entity","slug":"the-referential-reader-a-recurrent-entity","title":"The Referential Reader: A Recurrent Entity Network for Anaphora Resolution","date":"2019-02-05","arxiv_id":"1902.01541","repositories_listed":1,"syntology":null},{"url":"/paper/review-conversational-reading-comprehension","slug":"review-conversational-reading-comprehension","title":"Review Conversational Reading Comprehension","date":"2019-02-03","arxiv_id":"1902.00821","repositories_listed":1,"syntology":null},{"url":"/paper/a-generalized-language-model-in-tensor-space","slug":"a-generalized-language-model-in-tensor-space","title":"A Generalized Language Model in Tensor Space","date":"2019-01-31","arxiv_id":"1901.11167","repositories_listed":1,"syntology":null},{"url":"/paper/tensorized-embedding-layers-for-efficient","slug":"tensorized-embedding-layers-for-efficient","title":"Tensorized Embedding Layers for Efficient Model Compression","date":"2019-01-30","arxiv_id":"1901.10787","repositories_listed":1,"syntology":null},{"url":"/paper/latent-normalizing-flows-for-discrete","slug":"latent-normalizing-flows-for-discrete","title":"Latent Normalizing Flows for Discrete Sequences","date":"2019-01-29","arxiv_id":"1901.10548","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":1,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":2,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/latent-normalizing-flows-for-discrete#ran","syntology_url":"https://syntology.ai/paper/1901.10548","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1901.10548"}},"official":{"repos":["harvardnlp/TextFlow"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/fastgrnn-a-fast-accurate-stable-and-tiny","slug":"fastgrnn-a-fast-accurate-stable-and-tiny","title":"FastGRNN: A Fast, Accurate, Stable and Tiny Kilobyte Sized Gated Recurrent Neural Network","date":"2019-01-08","arxiv_id":"1901.02358","repositories_listed":1,"syntology":null},{"url":"/paper/team-papelo-transformer-networks-at-fever","slug":"team-papelo-transformer-networks-at-fever","title":"Team Papelo: Transformer Networks at FEVER","date":"2019-01-08","arxiv_id":"1901.02534","repositories_listed":1,"syntology":null},{"url":"/paper/transfer-learning-from-language-models-to","slug":"transfer-learning-from-language-models-to","title":"Transfer learning from language models to image caption generators: Better models may not transfer better","date":"2019-01-01","arxiv_id":"1901.01216","repositories_listed":1,"syntology":{"n":2,"n_ran":1,"n_constructed":0,"n_ran_checked":1,"n_instrument":0,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/transfer-learning-from-language-models-to#ran","syntology_url":"https://syntology.ai/paper/1901.01216","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1901.01216"}},"official":{"repos":["mtanti/mtanti-phd"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/exploring-weight-symmetry-in-deep-neural","slug":"exploring-weight-symmetry-in-deep-neural","title":"Exploring Weight Symmetry in Deep Neural Networks","date":"2018-12-28","arxiv_id":"1812.11027","repositories_listed":1,"syntology":{"n":8,"n_ran":5,"n_constructed":0,"n_ran_checked":4,"n_instrument":1,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":4,"n_pointer_only":1,"phrase":"5 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/exploring-weight-symmetry-in-deep-neural#ran","syntology_url":"https://syntology.ai/paper/1812.11027","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1812.11027"}},"official":{"repos":["hushell/deep-symmetry"],"state":"official (archive's flag): 5 ran","n_ran":5,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/writer-aware-cnn-for-parsimonious-hmm-based","slug":"writer-aware-cnn-for-parsimonious-hmm-based","title":"Writer-Aware CNN for Parsimonious HMM-Based Offline Handwritten Chinese Text Recognition","date":"2018-12-24","arxiv_id":"1812.09809","repositories_listed":1,"syntology":null},{"url":"/paper/what-is-one-grain-of-sand-in-the-desert","slug":"what-is-one-grain-of-sand-in-the-desert","title":"What Is One Grain of Sand in the Desert? Analyzing Individual Neurons in Deep NLP Models","date":"2018-12-21","arxiv_id":"1812.09355","repositories_listed":1,"syntology":null},{"url":"/paper/deep-networks-incorporating-spiking-neural","slug":"deep-networks-incorporating-spiking-neural","title":"Deep learning incorporating biologically-inspired neural dynamics","date":"2018-12-17","arxiv_id":"1812.07040","repositories_listed":1,"syntology":null},{"url":"/paper/recurrent-neural-networks-with-pre-trained","slug":"recurrent-neural-networks-with-pre-trained","title":"Recurrent Neural Networks with Pre-trained Language Model Embedding for Slot Filling Task","date":"2018-12-12","arxiv_id":"1812.05199","repositories_listed":1,"syntology":null},{"url":"/paper/practical-text-classification-with-large-pre","slug":"practical-text-classification-with-large-pre","title":"Practical Text Classification With Large Pre-Trained Language Models","date":"2018-12-04","arxiv_id":"1812.01207","repositories_listed":1,"syntology":null},{"url":"/paper/multi-level-multimodal-common-semantic-space","slug":"multi-level-multimodal-common-semantic-space","title":"Multi-level Multimodal Common Semantic Space for Image-Phrase Grounding","date":"2018-11-28","arxiv_id":"1811.11683","repositories_listed":1,"syntology":null},{"url":"/paper/alignment-analysis-of-sequential-segmentation","slug":"alignment-analysis-of-sequential-segmentation","title":"Alignment Analysis of Sequential Segmentation of Lexicons to Improve Automatic Cognate Detection","date":"2018-11-20","arxiv_id":"1811.08129","repositories_listed":1,"syntology":null},{"url":"/paper/unsupervised-transfer-learning-for-spoken","slug":"unsupervised-transfer-learning-for-spoken","title":"Unsupervised Transfer Learning for Spoken Language Understanding in Intelligent Agents","date":"2018-11-13","arxiv_id":"1811.05370","repositories_listed":1,"syntology":null},{"url":"/paper/effective-subtree-encoding-for-easy-first","slug":"effective-subtree-encoding-for-easy-first","title":"Effective Representation for Easy-First Dependency Parsing","date":"2018-11-08","arxiv_id":"1811.03511","repositories_listed":1,"syntology":null},{"url":"/paper/do-rnns-learn-human-like-abstract-word-order","slug":"do-rnns-learn-human-like-abstract-word-order","title":"Do RNNs learn human-like abstract word order preferences?","date":"2018-11-05","arxiv_id":"1811.01866","repositories_listed":1,"syntology":null},{"url":"/paper/mesh-tensorflow-deep-learning-for","slug":"mesh-tensorflow-deep-learning-for","title":"Mesh-TensorFlow: Deep Learning for Supercomputers","date":"2018-11-05","arxiv_id":"1811.02084","repositories_listed":1,"syntology":null},{"url":"/paper/sentence-encoders-on-stilts-supplementary","slug":"sentence-encoders-on-stilts-supplementary","title":"Sentence Encoders on STILTs: Supplementary Training on Intermediate Labeled-data Tasks","date":"2018-11-02","arxiv_id":"1811.01088","repositories_listed":1,"syntology":null},{"url":"/paper/juman-a-morphological-analysis-toolkit-for","slug":"juman-a-morphological-analysis-toolkit-for","title":"Juman++: A Morphological Analysis Toolkit for Scriptio Continua","date":"2018-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/on-the-end-to-end-solution-to-mandarin","slug":"on-the-end-to-end-solution-to-mandarin","title":"On the End-to-End Solution to Mandarin-English Code-switching Speech Recognition","date":"2018-11-01","arxiv_id":"1811.00241","repositories_listed":1,"syntology":null},{"url":"/paper/improving-machine-reading-comprehension-with","slug":"improving-machine-reading-comprehension-with","title":"Improving Machine Reading Comprehension with General Reading Strategies","date":"2018-10-31","arxiv_id":"1810.13441","repositories_listed":1,"syntology":null},{"url":"/paper/language-modeling-with-sparse-product-of","slug":"language-modeling-with-sparse-product-of","title":"Language Modeling with Sparse Product of Sememe Experts","date":"2018-10-29","arxiv_id":"1810.12387","repositories_listed":1,"syntology":null},{"url":"/paper/language-modeling-for-code-switching","slug":"language-modeling-for-code-switching","title":"Language Modeling for Code-Switching: Evaluation, Integration of Monolingual Data, and Discriminative Training","date":"2018-10-28","arxiv_id":"1810.11895","repositories_listed":1,"syntology":null},{"url":"/paper/evolutionary-stochastic-gradient-descent-for","slug":"evolutionary-stochastic-gradient-descent-for","title":"Evolutionary Stochastic Gradient Descent for Optimization of Deep Neural Networks","date":"2018-10-16","arxiv_id":"1810.06773","repositories_listed":1,"syntology":null},{"url":"/paper/trellis-networks-for-sequence-modeling","slug":"trellis-networks-for-sequence-modeling","title":"Trellis Networks for Sequence Modeling","date":"2018-10-15","arxiv_id":"1810.06682","repositories_listed":1,"syntology":{"n":8,"n_ran":7,"n_constructed":0,"n_ran_checked":7,"n_instrument":0,"n_unverified":1,"n_honours":2,"n_violates":0,"n_no_contract":5,"n_pointer_only":2,"phrase":"7 ran (of which 0 constructed an object rather than computing a result; 7 with no instrument failure: 2 honoured, 0 violated, 5 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/trellis-networks-for-sequence-modeling#ran","syntology_url":"https://syntology.ai/paper/1810.06682","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1810.06682"}},"official":{"repos":["locuslab/trellisnet"],"state":"official (archive's flag): 7 ran","n_ran":7,"n_constructed":0,"n_ran_no_instrument_failure":7,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/texttovec-deep-contextualized-neural","slug":"texttovec-deep-contextualized-neural","title":"textTOvec: Deep Contextualized Neural Autoregressive Topic Models of Language with Distributed Compositional Prior","date":"2018-10-09","arxiv_id":"1810.03947","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":0,"n_instrument":3,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 3 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/texttovec-deep-contextualized-neural#ran","syntology_url":"https://syntology.ai/paper/1810.03947","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1810.03947"}},"official":{"repos":["pgcool/textTOvec"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/unsupervised-neural-word-segmentation-for","slug":"unsupervised-neural-word-segmentation-for","title":"Unsupervised Neural Word Segmentation for Chinese via Segmental Language Modeling","date":"2018-10-07","arxiv_id":"1810.03167","repositories_listed":1,"syntology":null},{"url":"/paper/learning-compressed-transforms-with-low","slug":"learning-compressed-transforms-with-low","title":"Learning Compressed Transforms with Low Displacement Rank","date":"2018-10-04","arxiv_id":"1810.02309","repositories_listed":1,"syntology":{"n":21,"n_ran":10,"n_constructed":0,"n_ran_checked":10,"n_instrument":0,"n_unverified":11,"n_honours":0,"n_violates":0,"n_no_contract":10,"n_pointer_only":0,"phrase":"10 ran (of which 0 constructed an object rather than computing a result; 10 with no instrument failure: 0 honoured, 0 violated, 10 with no contract checked; 0 where Syntology's instrument failed) · 11 unverified","sample_list":"/paper/learning-compressed-transforms-with-low#ran","syntology_url":"https://syntology.ai/paper/1810.02309","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1810.02309"}},"official":{"repos":["HazyResearch/structured-nets"],"state":"official (archive's flag): 10 ran","n_ran":10,"n_constructed":0,"n_ran_no_instrument_failure":10,"n_unverified":11,"ran_from_kinds":["official"]}}},{"url":"/paper/a-hybrid-approach-to-automatic-corpus","slug":"a-hybrid-approach-to-automatic-corpus","title":"A Hybrid Approach to Automatic Corpus Generation for Chinese Spelling Check","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/diversity-promoting-gan-a-cross-entropy-based","slug":"diversity-promoting-gan-a-cross-entropy-based","title":"Diversity-Promoting GAN: A Cross-Entropy Based Generative Adversarial Network for Diversified Text Generation","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/prompsits-submission-to-wmt-2018-parallel","slug":"prompsits-submission-to-wmt-2018-parallel","title":"Prompsit's submission to WMT 2018 Parallel Corpus Filtering shared task","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/put-it-back-entity-typing-with-language-model","slug":"put-it-back-entity-typing-with-language-model","title":"Put It Back: Entity Typing with Language Model Enhancement","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/synthetic-data-made-to-order-the-case-of","slug":"synthetic-data-made-to-order-the-case-of","title":"Synthetic Data Made to Order: The Case of Parsing","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/learning-recurrent-binaryternary-weights","slug":"learning-recurrent-binaryternary-weights","title":"Learning Recurrent Binary/Ternary Weights","date":"2018-09-28","arxiv_id":"1809.11086","repositories_listed":1,"syntology":null},{"url":"/paper/controllable-neural-story-plot-generation-via","slug":"controllable-neural-story-plot-generation-via","title":"Controllable Neural Story Plot Generation via Reward Shaping","date":"2018-09-27","arxiv_id":"1809.10736","repositories_listed":1,"syntology":null},{"url":"/paper/document-informed-neural-autoregressive-topic","slug":"document-informed-neural-autoregressive-topic","title":"Document Informed Neural Autoregressive Topic Models with Distributional Prior","date":"2018-09-15","arxiv_id":"1809.06709","repositories_listed":1,"syntology":null},{"url":"/paper/rnns-as-psycholinguistic-subjects-syntactic","slug":"rnns-as-psycholinguistic-subjects-syntactic","title":"RNNs as psycholinguistic subjects: Syntactic state and grammatical dependency","date":"2018-09-05","arxiv_id":"1809.01329","repositories_listed":1,"syntology":null},{"url":"/paper/simple-fusion-return-of-the-language-model","slug":"simple-fusion-return-of-the-language-model","title":"Simple Fusion: Return of the Language Model","date":"2018-09-01","arxiv_id":"1809.00125","repositories_listed":1,"syntology":{"n":10,"n_ran":2,"n_constructed":0,"n_ran_checked":1,"n_instrument":1,"n_unverified":8,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 1 where Syntology's instrument failed) · 8 unverified","sample_list":"/paper/simple-fusion-return-of-the-language-model#ran","syntology_url":"https://syntology.ai/paper/1809.00125","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1809.00125"}},"official":null}},{"url":"/paper/spherical-latent-spaces-for-stable","slug":"spherical-latent-spaces-for-stable","title":"Spherical Latent Spaces for Stable Variational Autoencoders","date":"2018-08-31","arxiv_id":"1808.10805","repositories_listed":1,"syntology":{"n":6,"n_ran":6,"n_constructed":0,"n_ran_checked":6,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":6,"n_pointer_only":0,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 6 with no instrument failure: 0 honoured, 0 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/spherical-latent-spaces-for-stable#ran","syntology_url":"https://syntology.ai/paper/1808.10805","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1808.10805"}},"official":{"repos":["jiacheng-xu/vmf_vae_nlp"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":6,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/direct-output-connection-for-a-high-rank","slug":"direct-output-connection-for-a-high-rank","title":"Direct Output Connection for a High-Rank Language Model","date":"2018-08-30","arxiv_id":"1808.10143","repositories_listed":1,"syntology":null},{"url":"/paper/a-neural-model-of-adaptation-in-reading","slug":"a-neural-model-of-adaptation-in-reading","title":"A Neural Model of Adaptation in Reading","date":"2018-08-29","arxiv_id":"1808.09930","repositories_listed":1,"syntology":null},{"url":"/paper/grammar-induction-with-neural-language-models","slug":"grammar-induction-with-neural-language-models","title":"Grammar Induction with Neural Language Models: An Unusual Replication","date":"2018-08-29","arxiv_id":"1808.10000","repositories_listed":1,"syntology":null},{"url":"/paper/a-quantum-many-body-wave-function-inspired","slug":"a-quantum-many-body-wave-function-inspired","title":"A Quantum Many-body Wave Function Inspired Language Modeling Approach","date":"2018-08-28","arxiv_id":"1808.09891","repositories_listed":1,"syntology":null},{"url":"/paper/hierarchical-quantized-representations-for","slug":"hierarchical-quantized-representations-for","title":"Hierarchical Quantized Representations for Script Generation","date":"2018-08-28","arxiv_id":"1808.09542","repositories_listed":1,"syntology":null},{"url":"/paper/rational-recurrences","slug":"rational-recurrences","title":"Rational Recurrences","date":"2018-08-28","arxiv_id":"1808.09357","repositories_listed":1,"syntology":{"n":13,"n_ran":8,"n_constructed":0,"n_ran_checked":8,"n_instrument":0,"n_unverified":5,"n_honours":1,"n_violates":1,"n_no_contract":6,"n_pointer_only":1,"phrase":"8 ran (of which 0 constructed an object rather than computing a result; 8 with no instrument failure: 1 honoured, 1 violated, 6 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/rational-recurrences#ran","syntology_url":"https://syntology.ai/paper/1808.09357","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1808.09357"}},"official":{"repos":["Noahs-ARK/rational-recurrences"],"state":"official (archive's flag): 8 ran","n_ran":8,"n_constructed":0,"n_ran_no_instrument_failure":8,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/generating-text-through-adversarial-training","slug":"generating-text-through-adversarial-training","title":"Generating Text through Adversarial Training using Skip-Thought Vectors","date":"2018-08-27","arxiv_id":"1808.08703","repositories_listed":1,"syntology":null},{"url":"/paper/predefined-sparseness-in-recurrent-sequence","slug":"predefined-sparseness-in-recurrent-sequence","title":"Predefined Sparseness in Recurrent Sequence Models","date":"2018-08-27","arxiv_id":"1808.08720","repositories_listed":1,"syntology":null},{"url":"/paper/document-informed-neural-autoregressive-topic-1","slug":"document-informed-neural-autoregressive-topic-1","title":"Document Informed Neural Autoregressive Topic Models","date":"2018-08-11","arxiv_id":"1808.03793","repositories_listed":1,"syntology":null},{"url":"/paper/character-level-language-modeling-with-deeper","slug":"character-level-language-modeling-with-deeper","title":"Character-Level Language Modeling with Deeper Self-Attention","date":"2018-08-09","arxiv_id":"1808.04444","repositories_listed":1,"syntology":null},{"url":"/paper/large-scale-language-modeling-converging-on","slug":"large-scale-language-modeling-converging-on","title":"Large Scale Language Modeling: Converging on 40GB of Text in Four Hours","date":"2018-08-03","arxiv_id":"1808.01371","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":2,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/large-scale-language-modeling-converging-on#ran","syntology_url":"https://syntology.ai/paper/1808.01371","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1808.01371"}},"official":{"repos":["NVIDIA/sentiment-discovery"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/contextual-string-embeddings-for-sequence","slug":"contextual-string-embeddings-for-sequence","title":"Contextual String Embeddings for Sequence Labeling","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/on-device-neural-language-model-based-word","slug":"on-device-neural-language-model-based-word","title":"On-Device Neural Language Model Based Word Prediction","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/reproducing-and-regularizing-the-scrn-model","slug":"reproducing-and-regularizing-the-scrn-model","title":"Reproducing and Regularizing the SCRN Model","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/rnn-simulations-of-grammaticality-judgments","slug":"rnn-simulations-of-grammaticality-judgments","title":"RNN Simulations of Grammaticality Judgments on Long-distance Dependencies","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/a-comparison-of-techniques-for-language-model","slug":"a-comparison-of-techniques-for-language-model","title":"A Comparison of Techniques for Language Model Integration in Encoder-Decoder Speech Recognition","date":"2018-07-27","arxiv_id":"1807.10857","repositories_listed":1,"syntology":null},{"url":"/paper/bilingual-expert-can-find-translation-errors","slug":"bilingual-expert-can-find-translation-errors","title":"\"Bilingual Expert\" Can Find Translation Errors","date":"2018-07-25","arxiv_id":"1807.09433","repositories_listed":1,"syntology":null},{"url":"/paper/a-comparison-of-adaptation-techniques-and","slug":"a-comparison-of-adaptation-techniques-and","title":"A Comparison of Adaptation Techniques and Recurrent Neural Network Architectures","date":"2018-07-12","arxiv_id":"1807.06441","repositories_listed":1,"syntology":null},{"url":"/paper/deep-speare-a-joint-neural-model-of-poetic","slug":"deep-speare-a-joint-neural-model-of-poetic","title":"Deep-speare: A Joint Neural Model of Poetic Language, Meter and Rhyme","date":"2018-07-10","arxiv_id":"1807.03491","repositories_listed":1,"syntology":null},{"url":"/paper/improved-training-of-neural-trans-dimensional","slug":"improved-training-of-neural-trans-dimensional","title":"Improved training of neural trans-dimensional random field language models with dynamic noise-contrastive estimation","date":"2018-07-03","arxiv_id":"1807.00993","repositories_listed":1,"syntology":null},{"url":"/paper/baseline-a-library-for-rapid-modeling","slug":"baseline-a-library-for-rapid-modeling","title":"Baseline: A Library for Rapid Modeling, Experimentation and Development of Deep Learning Algorithms targeting NLP","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/document-modeling-with-external-attention-for","slug":"document-modeling-with-external-attention-for","title":"Document Modeling with External Attention for Sentence Extraction","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/subword-level-word-vector-representations-for","slug":"subword-level-word-vector-representations-for","title":"Subword-level Word Vector Representations for Korean","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/the-influence-of-context-on-sentence","slug":"the-influence-of-context-on-sentence","title":"The Influence of Context on Sentence Acceptability Judgements","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/word-error-rate-estimation-for-speech","slug":"word-error-rate-estimation-for-speech","title":"Word Error Rate Estimation for Speech Recognition: e-WER","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/handling-massive-n-gram-datasets-efficiently","slug":"handling-massive-n-gram-datasets-efficiently","title":"Handling Massive N-Gram Datasets Efficiently","date":"2018-06-25","arxiv_id":"1806.09447","repositories_listed":1,"syntology":null},{"url":"/paper/a-melody-conditioned-lyrics-language-model","slug":"a-melody-conditioned-lyrics-language-model","title":"A Melody-Conditioned Lyrics Language Model","date":"2018-06-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/neural-sign-language-translation","slug":"neural-sign-language-translation","title":"Neural Sign Language Translation","date":"2018-06-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/unsupervised-text-style-transfer-using","slug":"unsupervised-text-style-transfer-using","title":"Unsupervised Text Style Transfer using Language Models as Discriminators","date":"2018-05-30","arxiv_id":"1805.11749","repositories_listed":1,"syntology":null},{"url":"/paper/implicit-language-model-in-lstm-for-ocr","slug":"implicit-language-model-in-lstm-for-ocr","title":"Implicit Language Model in LSTM for OCR","date":"2018-05-23","arxiv_id":"1805.09441","repositories_listed":1,"syntology":null},{"url":"/paper/pushing-the-bounds-of-dropout","slug":"pushing-the-bounds-of-dropout","title":"Pushing the bounds of dropout","date":"2018-05-23","arxiv_id":"1805.09208","repositories_listed":1,"syntology":null},{"url":"/paper/a-simple-cache-model-for-image-recognition-1","slug":"a-simple-cache-model-for-image-recognition-1","title":"A Simple Cache Model for Image Recognition","date":"2018-05-21","arxiv_id":"1805.08709","repositories_listed":1,"syntology":null},{"url":"/paper/character-based-neural-networks-for-sentence","slug":"character-based-neural-networks-for-sentence","title":"Character-based Neural Networks for Sentence Pair Modeling","date":"2018-05-21","arxiv_id":"1805.08297","repositories_listed":1,"syntology":null},{"url":"/paper/numeracy-for-language-models-evaluating-and","slug":"numeracy-for-language-models-evaluating-and","title":"Numeracy for Language Models: Evaluating and Improving their Ability to Predict Numbers","date":"2018-05-21","arxiv_id":"1805.08154","repositories_listed":1,"syntology":null},{"url":"/paper/semstyle-learning-to-generate-stylised-image","slug":"semstyle-learning-to-generate-stylised-image","title":"SemStyle: Learning to Generate Stylised Image Captions using Unaligned Text","date":"2018-05-18","arxiv_id":"1805.07030","repositories_listed":1,"syntology":null},{"url":"/paper/a-context-based-approach-for-dialogue-act","slug":"a-context-based-approach-for-dialogue-act","title":"A Context-based Approach for Dialogue Act Recognition using Simple Recurrent Neural Networks","date":"2018-05-16","arxiv_id":"1805.06280","repositories_listed":1,"syntology":null},{"url":"/paper/polite-dialogue-generation-without-parallel","slug":"polite-dialogue-generation-without-parallel","title":"Polite Dialogue Generation Without Parallel Data","date":"2018-05-08","arxiv_id":"1805.03162","repositories_listed":1,"syntology":{"n":2,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":0,"n_honours":1,"n_violates":1,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 1 honoured, 1 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/polite-dialogue-generation-without-parallel#ran","syntology_url":"https://syntology.ai/paper/1805.03162","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1805.03162"}},"official":null}},{"url":"/paper/hierarchical-pointer-generator-memory-network","slug":"hierarchical-pointer-generator-memory-network","title":"Disentangling Language and Knowledge in Task-Oriented Dialogs","date":"2018-05-03","arxiv_id":"1805.01216","repositories_listed":1,"syntology":null},{"url":"/paper/evaluation-phonemic-transcription-of-low","slug":"evaluation-phonemic-transcription-of-low","title":"Evaluation Phonemic Transcription of Low-Resource Tonal Languages for Language Documentation","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/fonbund-a-library-for-combining-cross-lingual","slug":"fonbund-a-library-for-combining-cross-lingual","title":"FonBund: A Library for Combining Cross-lingual Phonological Segment Data","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/preparation-and-usage-of-xhosa","slug":"preparation-and-usage-of-xhosa","title":"Preparation and Usage of Xhosa Lexicographical Data for a Multilingual, Federated Environment","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/tf-lm-tensorflow-based-language-modeling","slug":"tf-lm-tensorflow-based-language-modeling","title":"TF-LM: TensorFlow-based Language Modeling Toolkit","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/syllable-based-sequence-to-sequence-speech","slug":"syllable-based-sequence-to-sequence-speech","title":"Syllable-Based Sequence-to-Sequence Speech Recognition with the Transformer in Mandarin Chinese","date":"2018-04-28","arxiv_id":"1804.10752","repositories_listed":1,"syntology":null},{"url":"/paper/spell-once-summon-anywhere-a-two-level-open","slug":"spell-once-summon-anywhere-a-two-level-open","title":"Spell Once, Summon Anywhere: A Two-Level Open-Vocabulary Language Model","date":"2018-04-23","arxiv_id":"1804.08205","repositories_listed":1,"syntology":{"n":5,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":3,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","sample_list":"/paper/spell-once-summon-anywhere-a-two-level-open#ran","syntology_url":"https://syntology.ai/paper/1804.08205","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1804.08205"}},"official":{"repos":["sjmielke/spell-once"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":3,"ran_from_kinds":["official"]}}},{"url":"/paper/efficient-contextualized-representation","slug":"efficient-contextualized-representation","title":"Efficient Contextualized Representation: Language Model Pruning for Sequence Labeling","date":"2018-04-20","arxiv_id":"1804.07827","repositories_listed":1,"syntology":null},{"url":"/paper/identifying-compromised-accounts-on-social","slug":"identifying-compromised-accounts-on-social","title":"Semantic Text Analysis for Detection of Compromised Accounts on Social Networks","date":"2018-04-19","arxiv_id":"1804.07247","repositories_listed":1,"syntology":null},{"url":"/paper/automated-evaluation-of-out-of-context-errors","slug":"automated-evaluation-of-out-of-context-errors","title":"Automated Evaluation of Out-of-Context Errors","date":"2018-03-23","arxiv_id":"1803.08983","repositories_listed":1,"syntology":null},{"url":"/paper/neural-lattice-language-models","slug":"neural-lattice-language-models","title":"Neural Lattice Language Models","date":"2018-03-13","arxiv_id":"1803.05071","repositories_listed":1,"syntology":null},{"url":"/paper/the-importance-of-being-recurrent-for","slug":"the-importance-of-being-recurrent-for","title":"The Importance of Being Recurrent for Modeling Hierarchical Structure","date":"2018-03-09","arxiv_id":"1803.03585","repositories_listed":1,"syntology":null},{"url":"/paper/deep-fsmn-for-large-vocabulary-continuous","slug":"deep-fsmn-for-large-vocabulary-continuous","title":"Deep-FSMN for Large Vocabulary Continuous Speech Recognition","date":"2018-03-04","arxiv_id":"1803.05030","repositories_listed":1,"syntology":null},{"url":"/paper/recurrent-neural-network-based-semantic","slug":"recurrent-neural-network-based-semantic","title":"Recurrent Neural Network-Based Semantic Variational Autoencoder for Sequence-to-Sequence Learning","date":"2018-02-09","arxiv_id":"1802.03238","repositories_listed":1,"syntology":null},{"url":"/paper/learning-from-past-mistakes-improving","slug":"learning-from-past-mistakes-improving","title":"Learning from Past Mistakes: Improving Automatic Speech Recognition Output via Noisy-Clean Phrase Context Modeling","date":"2018-02-07","arxiv_id":"1802.02607","repositories_listed":1,"syntology":null},{"url":"/paper/nested-lstms","slug":"nested-lstms","title":"Nested LSTMs","date":"2018-01-31","arxiv_id":"1801.10308","repositories_listed":1,"syntology":null},{"url":"/paper/strassennets-deep-learning-with-a","slug":"strassennets-deep-learning-with-a","title":"StrassenNets: Deep Learning with a Multiplication Budget","date":"2017-12-11","arxiv_id":"1712.03942","repositories_listed":1,"syntology":null},{"url":"/paper/contextualized-word-representations-for","slug":"contextualized-word-representations-for","title":"Contextualized Word Representations for Reading Comprehension","date":"2017-12-10","arxiv_id":"1712.03609","repositories_listed":1,"syntology":null}],"record_sha256":"7266f7d6da6991fe066aaf4e4a6dd29d8af81d9f5c0ebaaeb597105160568969","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}