{"url":"/dataset/locness-corpus","name":"WI-LOCNESS","full_name":"Cambridge English Write & Improve & LOCNESS","description_markdown":"WI-LOCNESS is part of the [Building Educational Applications 2019 Shared Task for Grammatical Error Correction](https://www.cl.cam.ac.uk/research/nl/bea2019st/). It consists of two datasets:\r\n\r\n- **LOCNESS**: is a corpus consisting of essays written by native English students. \r\n- **Cambridge English Write & Improve** (**W&I**): Write & Improve (Yannakoudakis et al., 2018) is an online web platform that assists non-native English students with their writing. Specifically, students from around the world submit letters, stories, articles and essays in response to various prompts, and the W&I system provides instant feedback. Since W&I went live in 2014, W&I annotators have manually annotated some of these submissions and assigned them a CEFR level.\r\n\r\nSource: [The BEA-2019 Shared Task on Grammatical Error Correction](/paper/the-bea-2019-shared-task-on-grammatical-error)\r\n\r\nImage source: [WI-LOCNESS](https://www.cl.cam.ac.uk/research/nl/bea2019st/#data)","description_withheld":null,"homepage":"https://www.cl.cam.ac.uk/research/nl/bea2019st/","introduced_date":null,"introduced_date_note":null,"introduced_by":{"paper":"/paper/the-bea-2019-shared-task-on-grammatical-error","title":"The BEA-2019 Shared Task on Grammatical Error Correction","first_author":"Christopher Bryant","url":null},"license":null,"modalities":[{"name":"Texts","url":"/datasets/modality/texts"}],"tasks":[{"name":"Grammatical Error Correction","url":"/task/grammatical-error-correction","datasets_with_task":"/datasets/task/grammatical-error-correction"},{"name":"Variational Inference","url":"/task/variational-inference","datasets_with_task":"/datasets/task/variational-inference"}],"languages":[{"name":"English","url":"/datasets/language/english"}],"variants":["BEA-2019 (test)","WI-LOCNESS"],"data_loaders":[{"repo":"https://github.com/huggingface/datasets","url":"https://huggingface.co/datasets/bea2019st/wi_locness","frameworks":["tf","pytorch","jax"]},{"repo":"https://github.com/huggingface/datasets","url":"https://huggingface.co/datasets/wi_locness","frameworks":["tf","pytorch","jax"]}],"num_papers_in_archive":29,"source":{"archive":"pwc-archive (Hugging Face), CC BY-SA 4.0","snapshot":"2025-07-28"},"benchmarks":[{"leaderboard":"/sota/grammatical-error-correction-on-bea-2019-test","task":"Grammatical Error Correction","dataset_variant":"BEA-2019 (test)","rows":19,"metrics":["F0.5"],"first_row_in_archive_order":{"model":"Majority-voting ensemble on best 7 models","paper":"/paper/pillars-of-grammatical-error-correction","metrics":{"F0.5":"81.4"},"code_links":[{"title":"grammarly/pillars-of-gec","url":"https://github.com/grammarly/pillars-of-gec"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"},{"leaderboard":"/sota/grammatical-error-correction-on-wi-locness-1","task":"Grammatical Error Correction","dataset_variant":"WI-LOCNESS","rows":1,"metrics":["F0.5"],"first_row_in_archive_order":{"model":"RedPenNet","paper":"/paper/redpennet-for-grammatical-error-correction","metrics":{"F0.5":"77.60"},"code_links":[{"title":"webspellchecker/unlp-2023-shared-task","url":"https://github.com/webspellchecker/unlp-2023-shared-task"}]},"note":"rows are the archive's own order at snapshot; nothing here re-ranks them"}],"papers_with_a_benchmark_row":[{"paper":"/paper/efficient-and-interpretable-grammatical-error","title":"Efficient and Interpretable Grammatical Error Correction with Mixture of Experts","date":"2024-10-30","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/pillars-of-grammatical-error-correction","title":"Pillars of Grammatical Error Correction: Comprehensive Inspection Of Contemporary Approaches In The Era of Large Language Models","date":"2024-04-23","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/unsupervised-grammatical-error-correction","title":"Unsupervised Grammatical Error Correction Rivaling Supervised Methods","date":"2023-12-06","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/system-combination-via-quality-estimation-for","title":"System Combination via Quality Estimation for Grammatical Error Correction","date":"2023-10-23","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/improving-seq2seq-grammatical-error","title":"Improving Seq2Seq Grammatical Error Correction via Decoding Interventions","date":"2023-10-23","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/redpennet-for-grammatical-error-correction","title":"RedPenNet for Grammatical Error Correction: Outputs to Tokens, Attentions to Spans","date":"2023-09-19","rows_on_this_dataset":2,"code_links":1,"syntology":null},{"paper":"/paper/frustratingly-easy-system-combination-for","title":"Frustratingly Easy System Combination for Grammatical Error Correction","date":"2022-07-01","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/ensembling-and-knowledge-distilling-of-large-1","title":"Ensembling and Knowledge Distilling of Large Sequence Taggers for Grammatical Error Correction","date":"2022-03-24","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/improved-grammatical-error-correction-by","title":"Improved grammatical error correction by ranking elementary edits","date":"2021-11-16","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/lm-critic-language-models-for-unsupervised","title":"LM-Critic: Language Models for Unsupervised Grammatical Error Correction","date":"2021-09-14","rows_on_this_dataset":1,"code_links":3,"syntology":null},{"paper":"/paper/neural-quality-estimation-with-multiple","title":"Neural Quality Estimation with Multiple Hypotheses for Grammatical Error Correction","date":"2021-05-10","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/gector-grammatical-error-correction-tag-not","title":"GECToR -- Grammatical Error Correction: Tag, Not Rewrite","date":"2020-05-26","rows_on_this_dataset":2,"code_links":3,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":12,"samples_ran":2,"samples_unverified":10,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/encoder-decoder-models-can-benefit-from-pre","title":"Encoder-Decoder Models Can Benefit from Pre-trained Masked Language Models in Grammatical Error Correction","date":"2020-05-03","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/an-empirical-study-of-incorporating-pseudo","title":"An Empirical Study of Incorporating Pseudo Data into Grammatical Error Correction","date":"2019-09-02","rows_on_this_dataset":1,"code_links":1,"syntology":{"read_at":"2026-09-24T18:15:14+00:00","samples_harvested":3,"samples_ran":3,"samples_unverified":0,"pointer_only_for_licence":0,"claim":"Per-sample execution on synthesized fixtures; not a correctness claim."}},{"paper":"/paper/the-laix-systems-in-the-bea-2019-gec-shared","title":"The LAIX Systems in the BEA-2019 GEC Shared Task","date":"2019-08-01","rows_on_this_dataset":1,"code_links":0,"syntology":null},{"paper":"/paper/neural-grammatical-error-correction-systems","title":"Neural Grammatical Error Correction Systems with Unsupervised Pre-training on Synthetic Data","date":"2019-08-01","rows_on_this_dataset":1,"code_links":1,"syntology":null},{"paper":"/paper/a-neural-grammatical-error-correction-system","title":"A Neural Grammatical Error Correction System Built On Better Pre-training and Sequential Transfer Learning","date":"2019-07-02","rows_on_this_dataset":1,"code_links":2,"syntology":null},{"paper":"/paper/learning-to-combine-grammatical-error","title":"Learning to combine Grammatical Error Corrections","date":"2019-06-10","rows_on_this_dataset":1,"code_links":1,"syntology":null}],"syntology_totals":{"read_at":"2026-09-24T18:15:14+00:00","papers_with_samples":2,"samples_harvested":15,"samples_ran":5,"samples_unverified":10,"pointer_only_for_licence":0,"papers_with_no_sample_that_ran":0,"note":"the per-paper counts above, summed; not a rate"},"papers_note":"The archive never published its papers-using-dataset list; these are papers with a leaderboard row on this dataset's benchmarks."}