{"about":{"non_affiliation":"Code with Papers and Syntology are not affiliated with, endorsed by, or sponsored by Papers with Code, Meta, or the pwc-archive mirror.","site":"https://codewithpapers.app","licence":"CC BY-SA 4.0","licence_url":"https://creativecommons.org/licenses/by-sa/4.0/legalcode","attribution":"https://codewithpapers.app/attribution","modified":"archive material modified by Syntology; see the attribution page","syntology":{"site":"https://syntology.ai","developers":"https://syntology.ai/developers","mcp":{"server":"https://syntology.ai/mcp","transport":"streamable-http","server_card":"https://syntology.ai/.well-known/mcp/server-card.json","auth":{"type":"trial token, no account","trial_token":"https://syntology.ai/api/oauth/trial/token","method":"POST","docs":"https://syntology.ai/developers"}},"have":"https://syntology.ai/api/graph/have?x=<method, arXiv id or title> (free, answers coverage only)","paper_base":"https://syntology.ai/paper/","atlas_base":"https://app.syntology.ai/?focus="},"machine_readable":[{"url":"https://codewithpapers.app/llms.txt","what":"the machine catalog: every machine-readable file, counted"},{"url":"https://codewithpapers.app/index/manifest.json","what":"paper-to-code index by arXiv id, with Syntology's counts"},{"url":"https://codewithpapers.app/search/manifest.json","what":"site search index (titles, authors) and its files"},{"url":"https://codewithpapers.app/download","what":"bulk files: Syntology's layer, described there"},{"url":"https://codewithpapers.app/build_manifest.json","what":"the build record: inputs, counts, exclusions, probes"}]},"url":"/task/part-of-speech-tagging/papers/2","list_of":"/task/part-of-speech-tagging","task":"Part-Of-Speech Tagging","archive":{"snapshot":"2025-07-28"},"key_notes":{"n_ran_checked":"legacy name, kept unchanged so existing readers do not break: it counts the samples that ran with no instrument failure (honoured, violated, and ran with no contract checked); it does not mean a contract was checked, and the pages print it as 'K with no instrument failure', not 'K checked'","n_constructed":"a sub-count of the samples that ran, never subtracted from them and never a failure: an executed sample whose run returned an instance of its own class (fixture_out_type equals the entry name): the run built an object and did not compute a result (Syntology's RAN record, counts.constructed)"},"syntology_read_at":"2026-09-28T10:30:06+00:00","order":"archive","order_definition":"repositories listed in the archive (most first), then date (newest first), then slug","page":2,"pages_in_order":10,"rows_per_page":100,"rows":[101,200],"of":990,"counts":{"archive_papers_tagged":990,"with_a_code_link":228,"where_syntology_ran_a_sample":28,"not_listed_spam_title":0,"listed":990,"listed_where_code_ran":28,"where_syntology_ran_a_sample_split":{"with_a_run_with_no_instrument_failure":24,"every_run_a_failure_of_syntologys_instrument":4,"listed_with_a_run_with_no_instrument_failure":24,"listed_every_run_a_failure_of_syntologys_instrument":4,"filter":{"states":["a run with no instrument failure","any run, instrument failures included"],"default":"a run with no instrument failure","note":"on the 'only where code ran' pages the default hides, in the browser, the rows where every run was a failure of Syntology's instrument; the second state shows them again. Rows are hidden, never re-ordered; these twins list every row"}},"definition":"distinct papers the archive tags; 'where Syntology ran a sample' counts papers with at least one harvested sample that ran, which is not a correctness claim"},"first_page":"/task/part-of-speech-tagging","prev":"/task/part-of-speech-tagging","next":"/task/part-of-speech-tagging/papers/3","papers":[{"url":"/paper/multilexnorm-a-shared-task-on-multilingual","slug":"multilexnorm-a-shared-task-on-multilingual","title":"MultiLexNorm: A Shared Task on Multilingual Lexical Normalization","date":"2021-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/test-harder-than-you-train-probing-with","slug":"test-harder-than-you-train-probing-with","title":"Test Harder than You Train: Probing with Extrapolation Splits","date":"2021-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/slovakbert-slovak-masked-language-model","slug":"slovakbert-slovak-masked-language-model","title":"SlovakBERT: Slovak Masked Language Model","date":"2021-09-30","arxiv_id":"2109.15254","repositories_listed":1,"syntology":null},{"url":"/paper/gernermed-an-open-german-medical-ner-model","slug":"gernermed-an-open-german-medical-ner-model","title":"GERNERMED -- An Open German Medical NER Model","date":"2021-09-24","arxiv_id":"2109.12104","repositories_listed":1,"syntology":null},{"url":"/paper/efficient-test-time-adapter-ensembling-for","slug":"efficient-test-time-adapter-ensembling-for","title":"Efficient Test Time Adapter Ensembling for Low-resource Language Varieties","date":"2021-09-10","arxiv_id":"2109.04877","repositories_listed":1,"syntology":null},{"url":"/paper/elit-emory-language-and-information-toolkit","slug":"elit-emory-language-and-information-toolkit","title":"ELIT: Emory Language and Information Toolkit","date":"2021-09-08","arxiv_id":"2109.03903","repositories_listed":1,"syntology":null},{"url":"/paper/an-ensemble-approach-for-annotating-source","slug":"an-ensemble-approach-for-annotating-source","title":"An Ensemble Approach for Annotating Source Code Identifiers with Part-of-speech Tags","date":"2021-09-01","arxiv_id":"2109.00629","repositories_listed":1,"syntology":null},{"url":"/paper/differences-in-chinese-and-western-tourists","slug":"differences-in-chinese-and-western-tourists","title":"Differences in Chinese and Western tourists faced with Japanese hospitality: A natural language processing approach","date":"2021-07-30","arxiv_id":"2107.14681","repositories_listed":1,"syntology":null},{"url":"/paper/specializing-multilingual-language-models-an","slug":"specializing-multilingual-language-models-an","title":"Specializing Multilingual Language Models: An Empirical Study","date":"2021-06-16","arxiv_id":"2106.09063","repositories_listed":1,"syntology":null},{"url":"/paper/minimax-and-neyman-pearson-meta-learning-for","slug":"minimax-and-neyman-pearson-meta-learning-for","title":"Minimax and Neyman-Pearson Meta-Learning for Outlier Languages","date":"2021-06-02","arxiv_id":"2106.01051","repositories_listed":1,"syntology":null},{"url":"/paper/weighted-training-for-cross-task-learning","slug":"weighted-training-for-cross-task-learning","title":"Weighted Training for Cross-Task Learning","date":"2021-05-28","arxiv_id":"2105.14095","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":2,"n_instrument":0,"n_unverified":1,"n_honours":2,"n_violates":0,"n_no_contract":0,"n_pointer_only":3,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 2 honoured, 0 violated, 0 with no contract checked; 0 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/weighted-training-for-cross-task-learning#ran","syntology_url":"https://syntology.ai/paper/2105.14095","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2105.14095"}},"official":{"repos":["HornHehhf/TAWT"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/the-interplay-between-morphological-typology","slug":"the-interplay-between-morphological-typology","title":"The interplay between language similarity and script on a novel multi-layer Algerian dialect corpus","date":"2021-05-16","arxiv_id":"2105.07400","repositories_listed":1,"syntology":null},{"url":"/paper/lexicon-enhanced-chinese-sequence-labelling","slug":"lexicon-enhanced-chinese-sequence-labelling","title":"Lexicon Enhanced Chinese Sequence Labeling Using BERT Adapter","date":"2021-05-15","arxiv_id":"2105.07148","repositories_listed":1,"syntology":null},{"url":"/paper/part-of-speech-tagging-of-swedish-texts-in","slug":"part-of-speech-tagging-of-swedish-texts-in","title":"Part-of-speech tagging of Swedish texts in the neural era","date":"2021-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/neural-sequence-segmentation-as-determining","slug":"neural-sequence-segmentation-as-determining","title":"Neural Sequence Segmentation as Determining the Leftmost Segments","date":"2021-04-15","arxiv_id":"2104.07217","repositories_listed":1,"syntology":{"n":8,"n_ran":6,"n_constructed":0,"n_ran_checked":2,"n_instrument":4,"n_unverified":2,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":8,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 4 where Syntology's instrument failed) · 2 unverified","sample_list":"/paper/neural-sequence-segmentation-as-determining#ran","syntology_url":"https://syntology.ai/paper/2104.07217","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2104.07217"}},"official":{"repos":["LeePleased/LeftmostSeg"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":2,"ran_from_kinds":["official"]}}},{"url":"/paper/frake-fusional-real-time-automatic-keyword","slug":"frake-fusional-real-time-automatic-keyword","title":"FRAKE: Fusional Real-time Automatic Keyword Extraction","date":"2021-04-10","arxiv_id":"2104.04830","repositories_listed":1,"syntology":null},{"url":"/paper/augmenting-part-of-speech-tagging-with","slug":"augmenting-part-of-speech-tagging-with","title":"Augmenting Part-of-speech Tagging with Syntactic Information for Vietnamese and Chinese","date":"2021-02-24","arxiv_id":"2102.12136","repositories_listed":1,"syntology":null},{"url":"/paper/trankit-a-light-weight-transformer-based","slug":"trankit-a-light-weight-transformer-based","title":"Trankit: A Light-Weight Transformer-based Toolkit for Multilingual Natural Language Processing","date":"2021-01-09","arxiv_id":"2101.03289","repositories_listed":1,"syntology":null},{"url":"/paper/phonlp-a-joint-multi-task-learning-model-for","slug":"phonlp-a-joint-multi-task-learning-model-for","title":"PhoNLP: A joint multi-task learning model for Vietnamese part-of-speech tagging, named entity recognition and dependency parsing","date":"2021-01-05","arxiv_id":"2101.01476","repositories_listed":1,"syntology":null},{"url":"/paper/assessing-emoji-use-in-modern-text-processing","slug":"assessing-emoji-use-in-modern-text-processing","title":"Assessing Emoji Use in Modern Text Processing Tools","date":"2021-01-02","arxiv_id":"2101.00430","repositories_listed":1,"syntology":null},{"url":"/paper/from-hero-to-z-eroe-a-benchmark-of-low-level","slug":"from-hero-to-z-eroe-a-benchmark-of-low-level","title":"From Hero to Z\\'eroe: A Benchmark of Low-Level Adversarial Attacks","date":"2020-12-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/increasing-learning-efficiency-of-self","slug":"increasing-learning-efficiency-of-self","title":"Increasing Learning Efficiency of Self-Attention Networks through Direct Position Interactions, Learnable Temperature, and Convoluted Attention","date":"2020-12-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/joint-chinese-word-segmentation-and-part-of-1","slug":"joint-chinese-word-segmentation-and-part-of-1","title":"Joint Chinese Word Segmentation and Part-of-speech Tagging via Multi-channel Attention of Character N-grams","date":"2020-12-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/re-framing-incremental-deep-language-models","slug":"re-framing-incremental-deep-language-models","title":"Re-framing Incremental Deep Language Models for Dialogue Processing with Multi-task Learning","date":"2020-11-13","arxiv_id":"2011.06754","repositories_listed":1,"syntology":null},{"url":"/paper/exploiting-cross-dialectal-gold-syntax-for","slug":"exploiting-cross-dialectal-gold-syntax-for","title":"Exploiting Cross-Dialectal Gold Syntax for Low-Resource Historical Languages: Towards a Generic Parser for Pre-Modern Slavic","date":"2020-11-12","arxiv_id":"2011.06467","repositories_listed":1,"syntology":null},{"url":"/paper/language-through-a-prism-a-spectral-approach","slug":"language-through-a-prism-a-spectral-approach","title":"Language Through a Prism: A Spectral Approach for Multiscale Language Representations","date":"2020-11-09","arxiv_id":"2011.04823","repositories_listed":1,"syntology":{"n":5,"n_ran":0,"n_constructed":0,"n_ran_checked":0,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"0 ran · 5 unverified","sample_list":"/paper/language-through-a-prism-a-spectral-approach#ran","syntology_url":"https://syntology.ai/paper/2011.04823","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2011.04823"}},"official":null}},{"url":"/paper/coding-textual-inputs-boosts-the-accuracy-of","slug":"coding-textual-inputs-boosts-the-accuracy-of","title":"Coding Textual Inputs Boosts the Accuracy of Neural Networks","date":"2020-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/effective-deep-learning-models-for-automatic","slug":"effective-deep-learning-models-for-automatic","title":"Effective Deep Learning Models for Automatic Diacritization of Arabic Text","date":"2020-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/from-hero-to-zeroe-a-benchmark-of-low-level","slug":"from-hero-to-zeroe-a-benchmark-of-low-level","title":"From Hero to Zéroe: A Benchmark of Low-Level Adversarial Attacks","date":"2020-10-12","arxiv_id":"2010.05648","repositories_listed":1,"syntology":{"n":3,"n_ran":2,"n_constructed":0,"n_ran_checked":0,"n_instrument":2,"n_unverified":1,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"2 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 1 unverified","sample_list":"/paper/from-hero-to-zeroe-a-benchmark-of-low-level#ran","syntology_url":"https://syntology.ai/paper/2010.05648","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2010.05648"}},"official":{"repos":["yannikbenz/zeroe"],"state":"official (archive's flag): 2 ran","n_ran":2,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":1,"ran_from_kinds":["official"]}}},{"url":"/paper/adversarial-attack-and-defense-of-structured","slug":"adversarial-attack-and-defense-of-structured","title":"Adversarial Attack and Defense of Structured Prediction Models","date":"2020-10-04","arxiv_id":"2010.01610","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":1,"n_ran_checked":1,"n_instrument":0,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":1,"n_pointer_only":1,"phrase":"1 ran (of which 1 constructed an object rather than computing a result; 1 with no instrument failure: 0 honoured, 0 violated, 1 with no contract checked; 0 where Syntology's instrument failed) · 0 unverified; the one sample that ran constructed an object rather than computing a result","sample_list":"/paper/adversarial-attack-and-defense-of-structured#ran","syntology_url":"https://syntology.ai/paper/2010.01610","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2010.01610"}},"official":{"repos":["WinnieHAN/structure_adv"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":1,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/n-ltp-a-open-source-neural-chinese-language","slug":"n-ltp-a-open-source-neural-chinese-language","title":"N-LTP: An Open-source Neural Language Technology Platform for Chinese","date":"2020-09-24","arxiv_id":"2009.11616","repositories_listed":1,"syntology":null},{"url":"/paper/latin-bert-a-contextual-language-model-for","slug":"latin-bert-a-contextual-language-model-for","title":"Latin BERT: A Contextual Language Model for Classical Philology","date":"2020-09-21","arxiv_id":"2009.10053","repositories_listed":1,"syntology":null},{"url":"/paper/persian-ezafe-recognition-using-transformers","slug":"persian-ezafe-recognition-using-transformers","title":"Persian Ezafe Recognition Using Transformers and Its Role in Part-Of-Speech Tagging","date":"2020-09-20","arxiv_id":"2009.09474","repositories_listed":1,"syntology":null},{"url":"/paper/fasthan-a-bert-based-joint-many-task-toolkit","slug":"fasthan-a-bert-based-joint-many-task-toolkit","title":"fastHan: A BERT-based Multi-Task Toolkit for Chinese NLP","date":"2020-09-18","arxiv_id":"2009.08633","repositories_listed":1,"syntology":null},{"url":"/paper/greek-bert-the-greeks-visiting-sesame-street","slug":"greek-bert-the-greeks-visiting-sesame-street","title":"GREEK-BERT: The Greeks visiting Sesame Street","date":"2020-08-27","arxiv_id":"2008.12014","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":2,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":2,"n_pointer_only":1,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 2 with no instrument failure: 0 honoured, 0 violated, 2 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/greek-bert-the-greeks-visiting-sesame-street#ran","syntology_url":"https://syntology.ai/paper/2008.12014","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2008.12014"}},"official":{"repos":["nlpaueb/greek-bert"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":2,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/reliable-part-of-speech-tagging-of-historical","slug":"reliable-part-of-speech-tagging-of-historical","title":"Reliable Part-of-Speech Tagging of Historical Corpora through Set-Valued Prediction","date":"2020-08-04","arxiv_id":"2008.01377","repositories_listed":1,"syntology":null},{"url":"/paper/playing-with-words-at-the-national-library-of","slug":"playing-with-words-at-the-national-library-of","title":"Playing with Words at the National Library of Sweden -- Making a Swedish BERT","date":"2020-07-03","arxiv_id":"2007.01658","repositories_listed":1,"syntology":null},{"url":"/paper/joint-chinese-word-segmentation-and-part-of","slug":"joint-chinese-word-segmentation-and-part-of","title":"Joint Chinese Word Segmentation and Part-of-speech Tagging via Two-way Attentions of Auto-analyzed Knowledge","date":"2020-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/are-all-languages-created-equal-in","slug":"are-all-languages-created-equal-in","title":"Are All Languages Created Equal in Multilingual BERT?","date":"2020-05-18","arxiv_id":"2005.09093","repositories_listed":1,"syntology":null},{"url":"/paper/exploring-and-predicting-transferability","slug":"exploring-and-predicting-transferability","title":"Exploring and Predicting Transferability across NLP Tasks","date":"2020-05-02","arxiv_id":"2005.00770","repositories_listed":1,"syntology":{"n":10,"n_ran":6,"n_constructed":0,"n_ran_checked":4,"n_instrument":2,"n_unverified":4,"n_honours":2,"n_violates":0,"n_no_contract":2,"n_pointer_only":4,"phrase":"6 ran (of which 0 constructed an object rather than computing a result; 4 with no instrument failure: 2 honoured, 0 violated, 2 with no contract checked; 2 where Syntology's instrument failed) · 4 unverified","sample_list":"/paper/exploring-and-predicting-transferability#ran","syntology_url":"https://syntology.ai/paper/2005.00770","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2005.00770"}},"official":{"repos":["tuvuumass/task-transferability"],"state":"official (archive's flag): 6 ran","n_ran":6,"n_constructed":0,"n_ran_no_instrument_failure":4,"n_unverified":4,"ran_from_kinds":["official"]}}},{"url":"/paper/a-focused-study-to-compare-arabic-pre","slug":"a-focused-study-to-compare-arabic-pre","title":"An Empirical Study of Pre-trained Transformers for Arabic Information Extraction","date":"2020-04-30","arxiv_id":"2004.14519","repositories_listed":1,"syntology":null},{"url":"/paper/a-cross-genre-ensemble-approach-to-robust","slug":"a-cross-genre-ensemble-approach-to-robust","title":"A Cross-Genre Ensemble Approach to Robust Reddit Part of Speech Tagging","date":"2020-04-29","arxiv_id":"2004.14312","repositories_listed":1,"syntology":null},{"url":"/paper/kvistur-2-0-a-bilstm-compound-splitter-for","slug":"kvistur-2-0-a-bilstm-compound-splitter-for","title":"Kvistur 2.0: a BiLSTM Compound Splitter for Icelandic","date":"2020-04-16","arxiv_id":"2004.07776","repositories_listed":1,"syntology":null},{"url":"/paper/phobert-pre-trained-language-models-for","slug":"phobert-pre-trained-language-models-for","title":"PhoBERT: Pre-trained language models for Vietnamese","date":"2020-03-02","arxiv_id":"2003.00744","repositories_listed":1,"syntology":null},{"url":"/paper/discontinuous-constituent-parsing-with","slug":"discontinuous-constituent-parsing-with","title":"Discontinuous Constituent Parsing with Pointer Networks","date":"2020-02-05","arxiv_id":"2002.01824","repositories_listed":1,"syntology":{"n":1,"n_ran":1,"n_constructed":0,"n_ran_checked":0,"n_instrument":1,"n_unverified":0,"n_honours":0,"n_violates":0,"n_no_contract":0,"n_pointer_only":1,"phrase":"1 ran (of which 0 constructed an object rather than computing a result; 0 with no instrument failure: 0 honoured, 0 violated, 0 with no contract checked; 1 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/discontinuous-constituent-parsing-with#ran","syntology_url":"https://syntology.ai/paper/2002.01824","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"2002.01824"}},"official":{"repos":["danifg/DiscoPointer"],"state":"official (archive's flag): 1 ran","n_ran":1,"n_constructed":0,"n_ran_no_instrument_failure":0,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/parameter-space-factorization-for-zero-shot","slug":"parameter-space-factorization-for-zero-shot","title":"Parameter Space Factorization for Zero-Shot Learning across Tasks and Languages","date":"2020-01-30","arxiv_id":"2001.11453","repositories_listed":1,"syntology":null},{"url":"/paper/sequence-labeling-approach-to-the-task-of","slug":"sequence-labeling-approach-to-the-task-of","title":"Sequence Labeling Approach to the Task of Sentence Boundary Detection","date":"2020-01-20","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/multilingual-is-not-enough-bert-for-finnish","slug":"multilingual-is-not-enough-bert-for-finnish","title":"Multilingual is not enough: BERT for Finnish","date":"2019-12-15","arxiv_id":"1912.07076","repositories_listed":1,"syntology":null},{"url":"/paper/morphological-tagging-and-lemmatization-of","slug":"morphological-tagging-and-lemmatization-of","title":"Morphological Tagging and Lemmatization of Albanian: A Manually Annotated Corpus and Neural Models","date":"2019-12-02","arxiv_id":"1912.00991","repositories_listed":1,"syntology":null},{"url":"/paper/language-agnostic-syllabification-with-neural","slug":"language-agnostic-syllabification-with-neural","title":"Language-Agnostic Syllabification with Neural Sequence Labeling","date":"2019-09-29","arxiv_id":"1909.13362","repositories_listed":1,"syntology":null},{"url":"/paper/from-english-to-code-switching-transfer","slug":"from-english-to-code-switching-transfer","title":"From English to Code-Switching: Transfer Learning with Strong Morphological Clues","date":"2019-09-11","arxiv_id":"1909.05158","repositories_listed":1,"syntology":null},{"url":"/paper/designing-and-interpreting-probes-with","slug":"designing-and-interpreting-probes-with","title":"Designing and Interpreting Probes with Control Tasks","date":"2019-09-08","arxiv_id":"1909.03368","repositories_listed":1,"syntology":null},{"url":"/paper/establishing-strong-baselines-for-the-new","slug":"establishing-strong-baselines-for-the-new","title":"Establishing Strong Baselines for the New Decade: Sequence Tagging, Syntactic and Semantic Parsing with BERT","date":"2019-08-14","arxiv_id":"1908.04943","repositories_listed":1,"syntology":null},{"url":"/paper/a-new-annotation-scheme-for-the-sejong-part","slug":"a-new-annotation-scheme-for-the-sejong-part","title":"A New Annotation Scheme for the Sejong Part-of-speech Tagged Corpus","date":"2019-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/reinforced-training-data-selection-for-domain","slug":"reinforced-training-data-selection-for-domain","title":"Reinforced Training Data Selection for Domain Adaptation","date":"2019-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/cross-lingual-syntactic-transfer-through","slug":"cross-lingual-syntactic-transfer-through","title":"Cross-Lingual Syntactic Transfer through Unsupervised Adaptation of Invertible Projections","date":"2019-06-06","arxiv_id":"1906.02656","repositories_listed":1,"syntology":null},{"url":"/paper/sequence-tagging-with-contextual-and-non","slug":"sequence-tagging-with-contextual-and-non","title":"Sequence Tagging with Contextual and Non-Contextual Subword Representations: A Multilingual Evaluation","date":"2019-06-04","arxiv_id":"1906.01569","repositories_listed":1,"syntology":null},{"url":"/paper/learning-task-specific-representation-for","slug":"learning-task-specific-representation-for","title":"Learning Task-specific Representation for Novel Words in Sequence Labeling","date":"2019-05-29","arxiv_id":"1905.12277","repositories_listed":1,"syntology":null},{"url":"/paper/a-grounded-unsupervised-universal-part-of","slug":"a-grounded-unsupervised-universal-part-of","title":"A Grounded Unsupervised Universal Part-of-Speech Tagger for Low-Resource Languages","date":"2019-04-10","arxiv_id":"1904.05426","repositories_listed":1,"syntology":null},{"url":"/paper/unsupervised-domain-adaptation-of","slug":"unsupervised-domain-adaptation-of","title":"Unsupervised Domain Adaptation of Contextualized Embeddings for Sequence Labeling","date":"2019-04-04","arxiv_id":"1904.02817","repositories_listed":1,"syntology":null},{"url":"/paper/vcwe-visual-character-enhanced-word","slug":"vcwe-visual-character-enhanced-word","title":"VCWE: Visual Character-Enhanced Word Embeddings","date":"2019-02-23","arxiv_id":"1902.08795","repositories_listed":1,"syntology":null},{"url":"/paper/girnet-interleaved-multi-task-recurrent-state","slug":"girnet-interleaved-multi-task-recurrent-state","title":"GIRNet: Interleaved Multi-Task Recurrent State Sequence Models","date":"2018-11-28","arxiv_id":"1811.11456","repositories_listed":1,"syntology":null},{"url":"/paper/juman-a-morphological-analysis-toolkit-for","slug":"juman-a-morphological-analysis-toolkit-for","title":"Juman++: A Morphological Analysis Toolkit for Scriptio Continua","date":"2018-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/a-hybrid-approach-to-automatic-corpus","slug":"a-hybrid-approach-to-automatic-corpus","title":"A Hybrid Approach to Automatic Corpus Generation for Chinese Spelling Check","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/evaluation-of-a-sequence-tagging-tool-for","slug":"evaluation-of-a-sequence-tagging-tool-for","title":"Evaluation of a Sequence Tagging Tool for Biomedical Texts","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/joint-learning-of-pos-and-dependencies-for","slug":"joint-learning-of-pos-and-dependencies-for","title":"Joint Learning of POS and Dependencies for Multilingual Universal Dependency Parsing","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/lemmatag-jointly-tagging-and-lemmatizing-for-1","slug":"lemmatag-jointly-tagging-and-lemmatizing-for-1","title":"LemmaTag: Jointly Tagging and Lemmatizing for Morphologically Rich Languages with BRNNs","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/speed-reading-learning-to-read-forbackward","slug":"speed-reading-learning-to-read-forbackward","title":"Speed Reading: Learning to Read ForBackward via Shuttle","date":"2018-10-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/toward-a-standardized-and-more-accurate","slug":"toward-a-standardized-and-more-accurate","title":"Toward a Standardized and More Accurate Indonesian Part-of-Speech Tagging","date":"2018-09-10","arxiv_id":"1809.03391","repositories_listed":1,"syntology":null},{"url":"/paper/towards-jointud-part-of-speech-tagging-and","slug":"towards-jointud-part-of-speech-tagging-and","title":"Towards JointUD: Part-of-speech Tagging and Lemmatization using Recurrent Neural Networks","date":"2018-09-10","arxiv_id":"1809.03211","repositories_listed":1,"syntology":null},{"url":"/paper/multi-source-domain-adaptation-with-mixture","slug":"multi-source-domain-adaptation-with-mixture","title":"Multi-Source Domain Adaptation with Mixture of Experts","date":"2018-09-07","arxiv_id":"1809.02256","repositories_listed":1,"syntology":{"n":3,"n_ran":3,"n_constructed":0,"n_ran_checked":1,"n_instrument":2,"n_unverified":0,"n_honours":1,"n_violates":0,"n_no_contract":0,"n_pointer_only":0,"phrase":"3 ran (of which 0 constructed an object rather than computing a result; 1 with no instrument failure: 1 honoured, 0 violated, 0 with no contract checked; 2 where Syntology's instrument failed) · 0 unverified","sample_list":"/paper/multi-source-domain-adaptation-with-mixture#ran","syntology_url":"https://syntology.ai/paper/1809.02256","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1809.02256"}},"official":{"repos":["jiangfeng1124/transfer"],"state":"official (archive's flag): 3 ran","n_ran":3,"n_constructed":0,"n_ran_no_instrument_failure":1,"n_unverified":0,"ran_from_kinds":["official"]}}},{"url":"/paper/distant-supervision-from-disparate-sources","slug":"distant-supervision-from-disparate-sources","title":"Distant Supervision from Disparate Sources for Low-Resource Part-of-Speech Tagging","date":"2018-08-29","arxiv_id":"1808.09733","repositories_listed":1,"syntology":null},{"url":"/paper/wisebe-window-based-sentence-boundary","slug":"wisebe-window-based-sentence-boundary","title":"WiSeBE: Window-based Sentence Boundary Evaluation","date":"2018-08-27","arxiv_id":"1808.08850","repositories_listed":1,"syntology":null},{"url":"/paper/contextual-string-embeddings-for-sequence","slug":"contextual-string-embeddings-for-sequence","title":"Contextual String Embeddings for Sequence Labeling","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/gold-standard-annotations-for-preposition-and","slug":"gold-standard-annotations-for-preposition-and","title":"Gold Standard Annotations for Preposition and Verb Sense with Semantic Role Labels in Adult-Child Interactions","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/learning-word-meta-embeddings-by-autoencoding","slug":"learning-word-meta-embeddings-by-autoencoding","title":"Learning Word Meta-Embeddings by Autoencoding","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/using-j-k-fold-cross-validation-to-reduce-1","slug":"using-j-k-fold-cross-validation-to-reduce-1","title":"Using J-K-fold Cross Validation To Reduce Variance When Tuning NLP Models","date":"2018-08-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/resource-size-matters-improving-neural-named","slug":"resource-size-matters-improving-neural-named","title":"Resource-Size matters: Improving Neural Named Entity Recognition with Optimized Large Corpora","date":"2018-07-26","arxiv_id":"1807.10675","repositories_listed":1,"syntology":null},{"url":"/paper/an-improved-neural-network-model-for-joint","slug":"an-improved-neural-network-model-for-joint","title":"An improved neural network model for joint POS tagging and dependency parsing","date":"2018-07-11","arxiv_id":"1807.03955","repositories_listed":1,"syntology":null},{"url":"/paper/a-multi-lingual-multi-task-architecture-for","slug":"a-multi-lingual-multi-task-architecture-for","title":"A Multi-lingual Multi-task Architecture for Low-resource Sequence Labeling","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/baseline-a-library-for-rapid-modeling","slug":"baseline-a-library-for-rapid-modeling","title":"Baseline: A Library for Rapid Modeling, Experimentation and Development of Deep Learning Algorithms targeting NLP","date":"2018-07-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/using-j-k-fold-cross-validation-to-reduce","slug":"using-j-k-fold-cross-validation-to-reduce","title":"Using J-K fold Cross Validation to Reduce Variance When Tuning NLP Models","date":"2018-06-19","arxiv_id":"1806.07139","repositories_listed":1,"syntology":null},{"url":"/paper/part-of-speech-tagging-on-an-endangered","slug":"part-of-speech-tagging-on-an-endangered","title":"Part-of-Speech Tagging on an Endangered Language: a Parallel Griko-Italian Resource","date":"2018-06-11","arxiv_id":"1806.03757","repositories_listed":1,"syntology":null},{"url":"/paper/gaussian-mixture-latent-vector-grammars","slug":"gaussian-mixture-latent-vector-grammars","title":"Gaussian Mixture Latent Vector Grammars","date":"2018-05-12","arxiv_id":"1805.04688","repositories_listed":1,"syntology":null},{"url":"/paper/chanot-an-intelligent-annotation-tool-for","slug":"chanot-an-intelligent-annotation-tool-for","title":"ChAnot: An Intelligent Annotation Tool for Indigenous and Highly Agglutinative Languages in Peru","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/creating-a-translation-matrix-of-the-bibleas","slug":"creating-a-translation-matrix-of-the-bibleas","title":"Creating a Translation Matrix of the Bible's Names Across 591 Languages","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/mgad-multilingual-generation-of-analogy","slug":"mgad-multilingual-generation-of-analogy","title":"MGAD: Multilingual Generation of Analogy Datasets","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/someweta-a-part-of-speech-tagger-for-german","slug":"someweta-a-part-of-speech-tagger-for-german","title":"SoMeWeTa: A Part-of-Speech Tagger for German Social Media and Web Texts","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/sudachi-a-japanese-tokenizer-for-business","slug":"sudachi-a-japanese-tokenizer-for-business","title":"Sudachi: a Japanese Tokenizer for Business","date":"2018-05-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/linguistically-informed-self-attention-for","slug":"linguistically-informed-self-attention-for","title":"Linguistically-Informed Self-Attention for Semantic Role Labeling","date":"2018-04-23","arxiv_id":"1804.08199","repositories_listed":1,"syntology":{"n":14,"n_ran":9,"n_constructed":0,"n_ran_checked":9,"n_instrument":0,"n_unverified":5,"n_honours":0,"n_violates":0,"n_no_contract":9,"n_pointer_only":0,"phrase":"9 ran (of which 0 constructed an object rather than computing a result; 9 with no instrument failure: 0 honoured, 0 violated, 9 with no contract checked; 0 where Syntology's instrument failed) · 5 unverified","sample_list":"/paper/linguistically-informed-self-attention-for#ran","syntology_url":"https://syntology.ai/paper/1804.08199","mcp":{"get_harvested_code_for_paper":{"arxiv_id":"1804.08199"}},"official":{"repos":["strubell/LISA"],"state":"official (archive's flag): 9 ran","n_ran":9,"n_constructed":0,"n_ran_no_instrument_failure":9,"n_unverified":5,"ran_from_kinds":["official"]}}},{"url":"/paper/parsing-tweets-into-universal-dependencies","slug":"parsing-tweets-into-universal-dependencies","title":"Parsing Tweets into Universal Dependencies","date":"2018-04-23","arxiv_id":"1804.08228","repositories_listed":1,"syntology":null},{"url":"/paper/evaluating-layers-of-representation-in-neural","slug":"evaluating-layers-of-representation-in-neural","title":"Evaluating Layers of Representation in Neural Machine Translation on Part-of-Speech and Semantic Tagging Tasks","date":"2018-01-23","arxiv_id":"1801.07772","repositories_listed":1,"syntology":null},{"url":"/paper/improving-the-accuracy-of-pre-trained-word","slug":"improving-the-accuracy-of-pre-trained-word","title":"Improving the Accuracy of Pre-trained Word Embeddings for Sentiment Analysis","date":"2017-11-23","arxiv_id":"1711.08609","repositories_listed":1,"syntology":null},{"url":"/paper/from-word-segmentation-to-pos-tagging-for","slug":"from-word-segmentation-to-pos-tagging-for","title":"From Word Segmentation to POS Tagging for Vietnamese","date":"2017-11-14","arxiv_id":"1711.04951","repositories_listed":1,"syntology":null},{"url":"/paper/robust-multilingual-part-of-speech-tagging","slug":"robust-multilingual-part-of-speech-tagging","title":"Robust Multilingual Part-of-Speech Tagging via Adversarial Training","date":"2017-11-14","arxiv_id":"1711.04903","repositories_listed":1,"syntology":null},{"url":"/paper/word-ordering-as-unsupervised-learning","slug":"word-ordering-as-unsupervised-learning","title":"Word Ordering as Unsupervised Learning Towards Syntactically Plausible Word Representations","date":"2017-11-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/semi-supervised-structured-prediction-with","slug":"semi-supervised-structured-prediction-with","title":"Semi-supervised Structured Prediction with Neural CRF Autoencoder","date":"2017-09-01","arxiv_id":null,"repositories_listed":1,"syntology":null},{"url":"/paper/nnvlp-a-neural-network-based-vietnamese","slug":"nnvlp-a-neural-network-based-vietnamese","title":"NNVLP: A Neural Network-Based Vietnamese Language Processing Toolkit","date":"2017-08-24","arxiv_id":"1708.07241","repositories_listed":1,"syntology":null},{"url":"/paper/learning-to-select-data-for-transfer-learning","slug":"learning-to-select-data-for-transfer-learning","title":"Learning to select data for transfer learning with Bayesian Optimization","date":"2017-07-17","arxiv_id":"1707.05246","repositories_listed":1,"syntology":null},{"url":"/paper/to-normalize-or-not-to-normalize-the-impact","slug":"to-normalize-or-not-to-normalize-the-impact","title":"To Normalize, or Not to Normalize: The Impact of Normalization on Part-of-Speech Tagging","date":"2017-07-17","arxiv_id":"1707.05116","repositories_listed":1,"syntology":null}],"record_sha256":"0af6d4955cd3f34d595969cb87fb6b4c3caa059ce65dd424752593459918511f","record_changed_at":"2026-09-28","record_changed_at_basis":"first_hashed"}