Building a Nasa Yuwe Language Corpus and Tagging with a Metaheuristic Approach. Sierra Martínez, L., M., Cobos, C., A., Corrales Muñoz, J., C., Rojas Curieux, T., Herrera-Viedma, E., & Peluffo-Ordóñez, D., H. Computación y Sistemas, 9, 2018.
Building a Nasa Yuwe Language Corpus and Tagging with a Metaheuristic Approach [link]Website  doi  abstract   bibtex   
Nasa Yuwe is the language of the Nasa indigenous community in Colombia. It is currently threatened with extinction. In this regard, a range of computer science solutions have been developed to the teaching and revitalization of the language. One of the most suitable approaches is the construction of a Part- Of-Speech Tagging (POST), which encourages the analysis and advanced processing of the language. Nevertheless, for Nasa Yuwe no tagged corpus exists, neither is there a POS Tagger and no related works have been reported. This paper therefore concentrates on building a linguistic corpus tagged for the Nasa Yuwe language and generating the first tagging application for Nasa Yuwe. The main results and findings are 1) the process of building the Nasa Yuwe corpus, 2) the tagsets and tagged sentences, as well as the statistics associated with the corpus, 3) results of two experiments to evaluate several POS Taggers (a Random tagger, three versions of HSTAGger, a tagger based on the harmony search metaheuristic, and three versions of a memetic algorithm GBHS Tagger, based on Global-Best Harmony Search (GBHS), Hill Climbing and an explicit Tabu memory, which obtained the best results in contrast with the other methods considered over the Nasa Yuwe language corpus.
@article{
 title = {Building a Nasa Yuwe Language Corpus and Tagging with a Metaheuristic Approach},
 type = {article},
 year = {2018},
 keywords = {Global-best harmony search,Harmony search,Hill climbing,Nasa Yuwe language,Part of speech tagger,Tabu memory,Tagged corpus},
 volume = {22},
 websites = {http://www.cys.cic.ipn.mx/ojs/index.php/CyS/article/view/3018},
 month = {9},
 day = {30},
 id = {5023710c-e1e1-3295-b86a-9ba69009ddc3},
 created = {2022-01-26T03:00:58.915Z},
 file_attached = {false},
 profile_id = {aba9653c-d139-3f95-aad8-969c487ed2f3},
 group_id = {b9022d50-068c-31b4-9174-ebfaaf9ee57b},
 last_modified = {2022-01-26T03:00:58.915Z},
 read = {false},
 starred = {false},
 authored = {false},
 confirmed = {true},
 hidden = {false},
 citation_key = {SierraMartinez2018},
 private_publication = {false},
 abstract = {Nasa Yuwe is the language of the Nasa indigenous community in Colombia. It is currently threatened with extinction. In this regard, a range of computer science solutions have been developed to the teaching and revitalization of the language. One of the most suitable approaches is the construction of a Part- Of-Speech Tagging (POST), which encourages the analysis and advanced processing of the language. Nevertheless, for Nasa Yuwe no tagged corpus exists, neither is there a POS Tagger and no related works have been reported. This paper therefore concentrates on building a linguistic corpus tagged for the Nasa Yuwe language and generating the first tagging application for Nasa Yuwe. The main results and findings are 1) the process of building the Nasa Yuwe corpus, 2) the tagsets and tagged sentences, as well as the statistics associated with the corpus, 3) results of two experiments to evaluate several POS Taggers (a Random tagger, three versions of HSTAGger, a tagger based on the harmony search metaheuristic, and three versions of a memetic algorithm GBHS Tagger, based on Global-Best Harmony Search (GBHS), Hill Climbing and an explicit Tabu memory, which obtained the best results in contrast with the other methods considered over the Nasa Yuwe language corpus.},
 bibtype = {article},
 author = {Sierra Martínez, Luz Marina and Cobos, Carlos Alberto and Corrales Muñoz, Juan Carlos and Rojas Curieux, Tulio and Herrera-Viedma, Enrique and Peluffo-Ordóñez, Diego Hernán},
 doi = {10.13053/cys-22-3-3018},
 journal = {Computación y Sistemas},
 number = {3}
}

Downloads: 0