{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/method/tanh-activation/papers/61","list_of":"/method/tanh-activation","method":"Tanh Activation","archive":{"snapshot":"2025-07-28"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"date (newest first), then slug","page":61,"pages_in_order":64,"rows_per_page":100,"rows":[6001,6100],"of":6333,"counts":{"archive_papers_tagged":6333,"with_a_code_link":2134,"where_syntology_ran_a_sample":386,"not_listed_spam_title":0,"listed":6333,"listed_where_code_ran":386,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":324,"every_run_a_failure_of_syntologys_instrument":62,"listed_with_a_run_with_no_instrument_failure":324,"listed_every_run_a_failure_of_syntologys_instrument":62,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/method/tanh-activation","prev":"/method/tanh-activation/papers/60","next":"/method/tanh-activation/papers/62","papers":[{"paper":null,"slug":"neural-decomposition-of-time-series-data-for","title":"Neural Decomposition of Time-Series Data for Effective Generalization","date":"2017-05-25","arxiv_id":"1705.09137","n_code_links":0,"syntology":null},{"paper":"/paper/deep-voice-2-multi-speaker-neural-text-to","slug":"deep-voice-2-multi-speaker-neural-text-to","title":"Deep Voice 2: Multi-Speaker Neural Text-to-Speech","date":"2017-05-24","arxiv_id":"1705.08947","n_code_links":1,"syntology":null},{"paper":null,"slug":"clinical-intervention-prediction-and","title":"Clinical Intervention Prediction and Understanding using Deep Networks","date":"2017-05-23","arxiv_id":"1705.08498","n_code_links":0,"syntology":null},{"paper":null,"slug":"efficiently-applying-attention-to-sequential","title":"Efficiently applying attention to sequential data with the Recurrent Discounted Attention unit","date":"2017-05-23","arxiv_id":"1705.08480","n_code_links":0,"syntology":null},{"paper":null,"slug":"techniques-for-visualizing-lstms-applied-to","title":"Techniques for visualizing LSTMs applied to electrocardiograms","date":"2017-05-23","arxiv_id":"1705.08153","n_code_links":0,"syntology":null},{"paper":null,"slug":"facial-expression-recognition-using-enhanced","title":"Facial Expression Recognition Using Enhanced Deep 3D Convolutional Neural Networks","date":"2017-05-22","arxiv_id":"1705.07871","n_code_links":0,"syntology":null},{"paper":null,"slug":"prediction-of-sea-surface-temperature-using","title":"Prediction of Sea Surface Temperature using Long Short-Term Memory","date":"2017-05-19","arxiv_id":"1705.06861","n_code_links":0,"syntology":null},{"paper":"/paper/parlai-a-dialog-research-software-platform","slug":"parlai-a-dialog-research-software-platform","title":"ParlAI: A Dialog Research Software Platform","date":"2017-05-18","arxiv_id":"1705.06476","n_code_links":23,"syntology":null},{"paper":null,"slug":"a-long-short-term-memory-recurrent-neural","title":"A Long Short-Term Memory Recurrent Neural Network Framework for Network Traffic Matrix Prediction","date":"2017-05-16","arxiv_id":"1705.05690","n_code_links":0,"syntology":null},{"paper":"/paper/shortfuse-biomedical-time-series","slug":"shortfuse-biomedical-time-series","title":"ShortFuse: Biomedical Time Series Representations in the Presence of Structured Information","date":"2017-05-13","arxiv_id":"1705.04790","n_code_links":4,"syntology":null},{"paper":null,"slug":"cham-action-recognition-using-convolutional","title":"CHAM: action recognition using convolutional hierarchical attention model","date":"2017-05-09","arxiv_id":"1705.03146","n_code_links":0,"syntology":null},{"paper":null,"slug":"deeptingle","title":"DeepTingle","date":"2017-05-09","arxiv_id":"1705.03557","n_code_links":0,"syntology":null},{"paper":null,"slug":"logical-parsing-from-natural-language-based","title":"Logical Parsing from Natural Language Based on a Neural Translation Model","date":"2017-05-09","arxiv_id":"1705.03389","n_code_links":0,"syntology":null},{"paper":"/paper/convolutional-sequence-to-sequence-learning","slug":"convolutional-sequence-to-sequence-learning","title":"Convolutional Sequence to Sequence Learning","date":"2017-05-08","arxiv_id":"1705.03122","n_code_links":37,"syntology":null},{"paper":null,"slug":"multi-resolution-lstm-for-long-term","title":"Multi Resolution LSTM For Long Term Prediction In Neural Activity Video","date":"2017-05-08","arxiv_id":"1705.02893","n_code_links":0,"syntology":null},{"paper":null,"slug":"context-aware-trajectory-prediction","title":"Context-Aware Trajectory Prediction","date":"2017-05-06","arxiv_id":"1705.02503","n_code_links":0,"syntology":null},{"paper":null,"slug":"max-pooling-loss-training-of-long-short-term","title":"Max-Pooling Loss Training of Long Short-Term Memory Networks for Small-Footprint Keyword Spotting","date":"2017-05-05","arxiv_id":"1705.02411","n_code_links":0,"syntology":null},{"paper":"/paper/am-i-done-predicting-action-progress-in","slug":"am-i-done-predicting-action-progress-in","title":"Am I Done? Predicting Action Progress in Videos","date":"2017-05-04","arxiv_id":"1705.01781","n_code_links":1,"syntology":null},{"paper":"/paper/recurrent-soft-attention-model-for-common","slug":"recurrent-soft-attention-model-for-common","title":"Recurrent Soft Attention Model for Common Object Recognition","date":"2017-05-04","arxiv_id":"1705.01921","n_code_links":1,"syntology":null},{"paper":"/paper/on-improving-deep-reinforcement-learning-for","slug":"on-improving-deep-reinforcement-learning-for","title":"On Improving Deep Reinforcement Learning for POMDPs","date":"2017-04-26","arxiv_id":"1704.07978","n_code_links":1,"syntology":null},{"paper":null,"slug":"joint-sequence-learning-and-cross-modality","title":"Joint Sequence Learning and Cross-Modality Convolution for 3D Biomedical Segmentation","date":"2017-04-25","arxiv_id":"1704.07754","n_code_links":0,"syntology":null},{"paper":null,"slug":"probabilistic-vehicle-trajectory-prediction","title":"Probabilistic Vehicle Trajectory Prediction over Occupancy Grid Map via Recurrent Neural Network","date":"2017-04-24","arxiv_id":"1704.07049","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-to-create-and-reuse-words-in-open","title":"Learning to Create and Reuse Words in Open-Vocabulary Neural Language Modeling","date":"2017-04-23","arxiv_id":"1704.06986","n_code_links":0,"syntology":null},{"paper":"/paper/learning-to-skim-text","slug":"learning-to-skim-text","title":"Learning to Skim Text","date":"2017-04-23","arxiv_id":"1704.06877","n_code_links":4,"syntology":null},{"paper":null,"slug":"affect-lm-a-neural-language-model-for","title":"Affect-LM: A Neural Language Model for Customizable Affective Text Generation","date":"2017-04-22","arxiv_id":"1704.06851","n_code_links":0,"syntology":null},{"paper":"/paper/diagonal-rnns-in-symbolic-music-modeling","slug":"diagonal-rnns-in-symbolic-music-modeling","title":"Diagonal RNNs in Symbolic Music Modeling","date":"2017-04-18","arxiv_id":"1704.05420","n_code_links":1,"syntology":null},{"paper":null,"slug":"land-cover-classification-via-multi-temporal","title":"Land Cover Classification via Multi-temporal Spatial Data by Recurrent Neural Networks","date":"2017-04-13","arxiv_id":"1704.04055","n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-set-based-metric-learning-for-video","title":"Attention-Set based Metric Learning for Video Face Recognition","date":"2017-04-12","arxiv_id":"1704.03805","n_code_links":0,"syntology":null},{"paper":null,"slug":"action-unit-detection-with-region-adaptation","title":"Action Unit Detection with Region Adaptation, Multi-labeling Learning and Optimal Temporal Fusing","date":"2017-04-10","arxiv_id":"1704.03067","n_code_links":0,"syntology":null},{"paper":null,"slug":"automated-curriculum-learning-for-neural","title":"Automated Curriculum Learning for Neural Networks","date":"2017-04-10","arxiv_id":"1704.03003","n_code_links":0,"syntology":null},{"paper":null,"slug":"learning-human-motion-models-for-long-term","title":"Learning Human Motion Models for Long-term Predictions","date":"2017-04-10","arxiv_id":"1704.02827","n_code_links":0,"syntology":null},{"paper":null,"slug":"conversation-modeling-on-reddit-using-a-graph","title":"Conversation Modeling on Reddit using a Graph-Structured LSTM","date":"2017-04-07","arxiv_id":"1704.02080","n_code_links":0,"syntology":null},{"paper":null,"slug":"two-stream-lstm-a-deep-fusion-framework-for","title":"Two Stream LSTM: A Deep Fusion Framework for Human Action Recognition","date":"2017-04-04","arxiv_id":"1704.01194","n_code_links":0,"syntology":null},{"paper":null,"slug":"attention-modeling-for-targeted-sentiment","title":"Attention Modeling for Targeted Sentiment","date":"2017-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"contextual-bidirectional-long-short-term","title":"Contextual Bidirectional Long Short-Term Memory Recurrent Neural Network Language Models: A Generative Approach to Sentiment Analysis","date":"2017-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/multilingual-lexicalized-constituency-parsing","slug":"multilingual-lexicalized-constituency-parsing","title":"Multilingual Lexicalized Constituency Parsing with Word-Level Auxiliary Tasks","date":"2017-04-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":null,"slug":"neural-networks-for-negation-cue-detection-in","title":"Neural Networks for Negation Cue Detection in Chinese","date":"2017-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"neural-temporal-relation-extraction","title":"Neural Temporal Relation Extraction","date":"2017-04-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":"/paper/factorization-tricks-for-lstm-networks","slug":"factorization-tricks-for-lstm-networks","title":"Factorization tricks for LSTM networks","date":"2017-03-31","arxiv_id":"1703.10722","n_code_links":2,"syntology":null},{"paper":null,"slug":"n-gram-language-modeling-using-recurrent","title":"N-gram Language Modeling using Recurrent Neural Network Estimation","date":"2017-03-31","arxiv_id":"1703.10724","n_code_links":0,"syntology":null},{"paper":"/paper/unpaired-image-to-image-translation-using","slug":"unpaired-image-to-image-translation-using","title":"Unpaired Image-to-Image Translation using Cycle-Consistent Adversarial Networks","date":"2017-03-30","arxiv_id":"1703.10593","n_code_links":190,"syntology":{"ran":14,"of":31,"n_ran_checked":9,"n_instrument":5,"unverified":17,"pointer_only":6,"phrase":"14 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 1 honoured, 0 violated, 8 with no contract checked; 5 where Syntology's instrument failed) · 17 unverified","official":{"repos":["junyanz/CycleGAN","junyanz/pytorch-CycleGAN-and-pix2pix"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":null,"slug":"pose-conditioned-spatio-temporal-attention","title":"Pose-conditioned Spatio-Temporal Attention for Human Action Recognition","date":"2017-03-29","arxiv_id":"1703.10106","n_code_links":0,"syntology":null},{"paper":"/paper/tacotron-towards-end-to-end-speech-synthesis","slug":"tacotron-towards-end-to-end-speech-synthesis","title":"Tacotron: Towards End-to-End Speech Synthesis","date":"2017-03-29","arxiv_id":"1703.10135","n_code_links":30,"syntology":{"ran":16,"of":25,"n_ran_checked":13,"n_instrument":3,"unverified":9,"pointer_only":6,"phrase":"16 ran (of which 0 constructed an object rather than computing a result; 13 with no instrument failure: 3 honoured, 1 violated, 9 with no contract checked; 3 where Syntology's instrument failed) · 9 unverified","official":null}},{"paper":null,"slug":"collective-anomaly-detection-based-on-long","title":"Collective Anomaly Detection based on Long Short Term Memory Recurrent Neural Network","date":"2017-03-28","arxiv_id":"1703.09752","n_code_links":0,"syntology":null},{"paper":null,"slug":"ensembles-of-deep-lstm-learners-for-activity","title":"Ensembles of Deep LSTM Learners for Activity Recognition using Wearables","date":"2017-03-28","arxiv_id":"1703.09370","n_code_links":0,"syntology":null},{"paper":"/paper/view-adaptive-recurrent-neural-networks-for","slug":"view-adaptive-recurrent-neural-networks-for","title":"View Adaptive Recurrent Neural Networks for High Performance Human Action Recognition from Skeleton Data","date":"2017-03-24","arxiv_id":"1703.08274","n_code_links":1,"syntology":null},{"paper":"/paper/recurrent-multimodal-interaction-for","slug":"recurrent-multimodal-interaction-for","title":"Recurrent Multimodal Interaction for Referring Image Segmentation","date":"2017-03-23","arxiv_id":"1703.07939","n_code_links":1,"syntology":null},{"paper":null,"slug":"deep-lstm-for-large-vocabulary-continuous","title":"Deep LSTM for Large Vocabulary Continuous Speech Recognition","date":"2017-03-21","arxiv_id":"1703.07090","n_code_links":0,"syntology":null},{"paper":"/paper/dance-dance-convolution","slug":"dance-dance-convolution","title":"Dance Dance Convolution","date":"2017-03-20","arxiv_id":"1703.06891","n_code_links":1,"syntology":null},{"paper":"/paper/encoding-sentences-with-graph-convolutional","slug":"encoding-sentences-with-graph-convolutional","title":"Encoding Sentences with Graph Convolutional Networks for Semantic Role Labeling","date":"2017-03-14","arxiv_id":"1703.04826","n_code_links":2,"syntology":{"ran":4,"of":7,"n_ran_checked":4,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"4 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 3 unverified","official":{"repos":["diegma/neural-dep-srl"],"state":"official (archive's flag): 4 ran","n_ran":4,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":3,"ran_from_kinds":["official"]}}},{"paper":null,"slug":"neural-graph-machines-learning-neural","title":"Neural Graph Machines: Learning Neural Networks Using Graphs","date":"2017-03-14","arxiv_id":"1703.04818","n_code_links":0,"syntology":null},{"paper":null,"slug":"a-hierarchical-framework-of-cloud-resource","title":"A Hierarchical Framework of Cloud Resource Allocation and Power Management Using Deep Reinforcement Learning","date":"2017-03-13","arxiv_id":"1703.04221","n_code_links":0,"syntology":null},{"paper":"/paper/dragnn-a-transition-based-framework-for","slug":"dragnn-a-transition-based-framework-for","title":"DRAGNN: A Transition-based Framework for Dynamically Connected Neural Networks","date":"2017-03-13","arxiv_id":"1703.04474","n_code_links":1,"syntology":null},{"paper":null,"slug":"story-cloze-ending-selection-baselines-and","title":"Story Cloze Ending Selection Baselines and Data Examination","date":"2017-03-13","arxiv_id":"1703.04330","n_code_links":0,"syntology":null},{"paper":null,"slug":"interpretable-structure-evolving-lstm","title":"Interpretable Structure-Evolving LSTM","date":"2017-03-08","arxiv_id":"1703.03055","n_code_links":0,"syntology":null},{"paper":"/paper/english-conversational-telephone-speech","slug":"english-conversational-telephone-speech","title":"English Conversational Telephone Speech Recognition by Humans and Machines","date":"2017-03-06","arxiv_id":"1703.02136","n_code_links":0,"syntology":null},{"paper":"/paper/machine-learning-on-sequential-data-using-a","slug":"machine-learning-on-sequential-data-using-a","title":"Machine Learning on Sequential Data Using a Recurrent Weighted Average","date":"2017-03-03","arxiv_id":"1703.01253","n_code_links":2,"syntology":{"ran":6,"of":6,"n_ran_checked":1,"n_instrument":5,"unverified":0,"pointer_only":6,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 5 where Syntology's instrument failed) · 0 unverified","official":{"repos":["jostmey/rwa"],"state":"community repositories only","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["listed"]}}},{"paper":"/paper/on-geometric-features-for-skeleton-based","slug":"on-geometric-features-for-skeleton-based","title":"On Geometric Features for Skeleton-Based Action Recognition using Multilayer LSTM Networks","date":"2017-03-01","arxiv_id":null,"n_code_links":1,"syntology":null},{"paper":"/paper/the-statistical-recurrent-unit","slug":"the-statistical-recurrent-unit","title":"The Statistical Recurrent Unit","date":"2017-03-01","arxiv_id":"1703.00381","n_code_links":2,"syntology":null},{"paper":"/paper/improved-variational-autoencoders-for-text","slug":"improved-variational-autoencoders-for-text","title":"Improved Variational Autoencoders for Text Modeling using Dilated Convolutions","date":"2017-02-27","arxiv_id":"1702.08139","n_code_links":3,"syntology":null},{"paper":"/paper/neural-map-structured-memory-for-deep","slug":"neural-map-structured-memory-for-deep","title":"Neural Map: Structured Memory for Deep Reinforcement Learning","date":"2017-02-27","arxiv_id":"1702.08360","n_code_links":1,"syntology":null},{"paper":null,"slug":"analyzing-and-exploiting-narx-recurrent","title":"Analyzing and Exploiting NARX Recurrent Neural Networks for Long-Term Dependencies","date":"2017-02-24","arxiv_id":"1702.07805","n_code_links":0,"syntology":null},{"paper":null,"slug":"how-hard-is-it-to-cross-the-room-training","title":"How hard is it to cross the room? -- Training (Recurrent) Neural Networks to steer a UAV","date":"2017-02-24","arxiv_id":"1702.07600","n_code_links":0,"syntology":null},{"paper":null,"slug":"task-driven-visual-saliency-and-attention","title":"Task-driven Visual Saliency and Attention-based Visual Question Answering","date":"2017-02-22","arxiv_id":"1702.06700","n_code_links":0,"syntology":null},{"paper":null,"slug":"progressively-diffused-networks-for-semantic","title":"Progressively Diffused Networks for Semantic Image Segmentation","date":"2017-02-20","arxiv_id":"1702.05839","n_code_links":0,"syntology":null},{"paper":null,"slug":"survey-of-reasoning-using-neural-networks","title":"Survey of reasoning using Neural networks","date":"2017-02-14","arxiv_id":"1702.06186","n_code_links":0,"syntology":null},{"paper":"/paper/bilateral-multi-perspective-matching-for","slug":"bilateral-multi-perspective-matching-for","title":"Bilateral Multi-Perspective Matching for Natural Language Sentences","date":"2017-02-13","arxiv_id":"1702.03814","n_code_links":10,"syntology":{"ran":1,"of":4,"n_ran_checked":0,"n_instrument":1,"unverified":3,"pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 3 unverified","official":null}},{"paper":null,"slug":"parallel-long-short-term-memory-for-multi","title":"Parallel Long Short-Term Memory for Multi-stream Classification","date":"2017-02-11","arxiv_id":"1702.03402","n_code_links":0,"syntology":null},{"paper":null,"slug":"automatic-rule-extraction-from-long-short","title":"Automatic Rule Extraction from Long Short Term Memory Networks","date":"2017-02-08","arxiv_id":"1702.02540","n_code_links":0,"syntology":null},{"paper":"/paper/a-knowledge-grounded-neural-conversation","slug":"a-knowledge-grounded-neural-conversation","title":"A Knowledge-Grounded Neural Conversation Model","date":"2017-02-07","arxiv_id":"1702.01932","n_code_links":2,"syntology":null},{"paper":null,"slug":"recurrent-neural-networks-for-anomaly","title":"Recurrent Neural Networks for anomaly detection in the Post-Mortem time series of LHC superconducting magnets","date":"2017-02-02","arxiv_id":"1702.00833","n_code_links":0,"syntology":null},{"paper":null,"slug":"sensegen-a-deep-learning-architecture-for","title":"SenseGen: A Deep Learning Architecture for Synthetic Sensor Data Generation","date":"2017-01-31","arxiv_id":"1701.08886","n_code_links":0,"syntology":null},{"paper":"/paper/drug-drug-interaction-extraction-from","slug":"drug-drug-interaction-extraction-from","title":"Drug-Drug Interaction Extraction from Biomedical Text Using Long Short Term Memory Network","date":"2017-01-28","arxiv_id":"1701.08303","n_code_links":1,"syntology":null},{"paper":"/paper/outrageously-large-neural-networks-the","slug":"outrageously-large-neural-networks-the","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","date":"2017-01-23","arxiv_id":"1701.06538","n_code_links":4,"syntology":{"ran":4,"of":6,"n_ran_checked":4,"n_instrument":0,"unverified":2,"pointer_only":6,"phrase":"4 ran (of which 4 constructed an object rather than computing a result; 4 with no instrument failure: 0 honoured, 0 violated, 4 with no contract checked; 0 where Syntology's instrument failed) · 2 unverified; every one of the 4 samples that ran constructed an object rather than computing a result","official":null}},{"paper":"/paper/person-re-identification-via-recurrent","slug":"person-re-identification-via-recurrent","title":"Person Re-Identification via Recurrent Feature Aggregation","date":"2017-01-23","arxiv_id":"1701.06351","n_code_links":1,"syntology":null},{"paper":null,"slug":"end-to-end-visual-speech-recognition-with","title":"End-To-End Visual Speech Recognition With LSTMs","date":"2017-01-20","arxiv_id":"1701.05847","n_code_links":0,"syntology":null},{"paper":null,"slug":"auxiliary-multimodal-lstm-for-audio-visual","title":"Auxiliary Multimodal LSTM for Audio-visual Speech Recognition and Lipreading","date":"2017-01-16","arxiv_id":"1701.04224","n_code_links":0,"syntology":null},{"paper":"/paper/simplified-gating-in-long-short-term-memory","slug":"simplified-gating-in-long-short-term-memory","title":"Simplified Gating in Long Short-term Memory (LSTM) Recurrent Neural Networks","date":"2017-01-12","arxiv_id":"1701.03441","n_code_links":1,"syntology":null},{"paper":"/paper/a-simple-and-accurate-syntax-agnostic-neural","slug":"a-simple-and-accurate-syntax-agnostic-neural","title":"A Simple and Accurate Syntax-Agnostic Neural Model for Dependency-based Semantic Role Labeling","date":"2017-01-10","arxiv_id":"1701.02593","n_code_links":2,"syntology":null},{"paper":"/paper/residual-lstm-design-of-a-deep-recurrent","slug":"residual-lstm-design-of-a-deep-recurrent","title":"Residual LSTM: Design of a Deep Recurrent Architecture for Distant Speech Recognition","date":"2017-01-10","arxiv_id":"1701.03360","n_code_links":3,"syntology":null},{"paper":null,"slug":"neural-probabilistic-model-for-non-projective","title":"Neural Probabilistic Model for Non-projective MST Parsing","date":"2017-01-04","arxiv_id":"1701.00874","n_code_links":0,"syntology":null},{"paper":null,"slug":"shortcut-sequence-tagging","title":"Shortcut Sequence Tagging","date":"2017-01-03","arxiv_id":"1701.00576","n_code_links":0,"syntology":null},{"paper":null,"slug":"head-lexicalized-bidirectional-tree-lstms","title":"Head-Lexicalized Bidirectional Tree LSTMs","date":"2017-01-01","arxiv_id":null,"n_code_links":0,"syntology":null},{"paper":null,"slug":"p-dla-a-predictive-system-model-for-onshore","title":"p-DLA: A Predictive System Model for Onshore Oil and Gas Pipeline Dataset Classification and Monitoring - Part 1","date":"2016-12-31","arxiv_id":"1701.00040","n_code_links":0,"syntology":null},{"paper":"/paper/the-neural-hawkes-process-a-neurally-self","slug":"the-neural-hawkes-process-a-neurally-self","title":"The Neural Hawkes Process: A Neurally Self-Modulating Multivariate Point Process","date":"2016-12-29","arxiv_id":"1612.09328","n_code_links":9,"syntology":null},{"paper":null,"slug":"heres-my-point-joint-pointer-architecture-for","title":"Here's My Point: Joint Pointer Architecture for Argument Mining","date":"2016-12-28","arxiv_id":"1612.08994","n_code_links":0,"syntology":null},{"paper":null,"slug":"robust-lstm-autoencoders-for-face-de","title":"Robust LSTM-Autoencoders for Face De-Occlusion in the Wild","date":"2016-12-27","arxiv_id":"1612.08534","n_code_links":0,"syntology":null},{"paper":null,"slug":"multivariate-industrial-time-series-with","title":"Multivariate Industrial Time Series with Cyber-Attack Simulation: Fault Detection Using an LSTM-based Predictive Data Model","date":"2016-12-20","arxiv_id":"1612.06676","n_code_links":0,"syntology":null},{"paper":"/paper/span-based-constituency-parsing-with-a","slug":"span-based-constituency-parsing-with-a","title":"Span-Based Constituency Parsing with a Structure-Label System and Provably Optimal Dynamic Oracles","date":"2016-12-20","arxiv_id":"1612.06475","n_code_links":1,"syntology":null},{"paper":null,"slug":"transition-based-parsing-with-context","title":"Transition-based Parsing with Context Enhancement and Future Reward Reranking","date":"2016-12-15","arxiv_id":"1612.05131","n_code_links":0,"syntology":null},{"paper":null,"slug":"real-time-interactive-sequence-generation-and","title":"Real-time interactive sequence generation and control with Recurrent Neural Network ensembles","date":"2016-12-14","arxiv_id":"1612.04687","n_code_links":0,"syntology":null},{"paper":"/paper/improving-neural-language-models-with-a","slug":"improving-neural-language-models-with-a","title":"Improving Neural Language Models with a Continuous Cache","date":"2016-12-13","arxiv_id":"1612.04426","n_code_links":14,"syntology":null},{"paper":"/paper/multi-perspective-context-matching-for","slug":"multi-perspective-context-matching-for","title":"Multi-Perspective Context Matching for Machine Comprehension","date":"2016-12-13","arxiv_id":"1612.04211","n_code_links":1,"syntology":null},{"paper":null,"slug":"a-unit-selection-methodology-for-music","title":"A Unit Selection Methodology for Music Generation Using Deep Neural Networks","date":"2016-12-12","arxiv_id":"1612.03789","n_code_links":0,"syntology":null},{"paper":null,"slug":"empirical-evaluation-of-a-new-approach-to","title":"Empirical Evaluation of A New Approach to Simplifying Long Short-term Memory (LSTM)","date":"2016-12-12","arxiv_id":"1612.03707","n_code_links":0,"syntology":null},{"paper":"/paper/tracking-the-world-state-with-recurrent","slug":"tracking-the-world-state-with-recurrent","title":"Tracking the World State with Recurrent Entity Networks","date":"2016-12-12","arxiv_id":"1612.03969","n_code_links":5,"syntology":{"ran":0,"of":3,"n_ran_checked":0,"n_instrument":0,"unverified":3,"pointer_only":0,"phrase":"0 ran · 3 unverified","official":{"repos":["facebook/MemNN"],"state":"official: not harvested","n_ran":0,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":[]}}},{"paper":null,"slug":"towards-better-decoding-and-language-model","title":"Towards better decoding and language model integration in sequence to sequence models","date":"2016-12-08","arxiv_id":"1612.02695","n_code_links":0,"syntology":null},{"paper":null,"slug":"neural-turing-machines-convergence-of-copy","title":"Neural Turing Machines: Convergence of Copy Tasks","date":"2016-12-07","arxiv_id":"1612.02336","n_code_links":0,"syntology":null},{"paper":null,"slug":"short-term-traffic-flow-forecasting-with","title":"Short-term traffic flow forecasting with spatial-temporal correlation in a hybrid deep learning framework","date":"2016-12-03","arxiv_id":"1612.01022","n_code_links":0,"syntology":null},{"paper":"/paper/shift-reduce-constituent-parsing-with-neural","slug":"shift-reduce-constituent-parsing-with-neural","title":"Shift-Reduce Constituent Parsing with Neural Lookahead Features","date":"2016-12-02","arxiv_id":"1612.00567","n_code_links":1,"syntology":null}],"record_sha256":"2c3634ade656a7605a3d323d29b3d40cf9692cd7c1171eaf2e0fd03237642c8f","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}