@inproceedings{cristea-etal-2020-cobiliro,
title = "{C}o{B}i{L}i{R}o: A Research Platform for Bimodal Corpora",
author = "Cristea, Dan and
Pistol, Ionuț and
Boghiu, Șerban and
Bibiri, Anca-Diana and
G{\^i}fu, Daniela and
Scutelnicu, Andrei and
Onofrei, Mihaela and
Trandab{\u{a}}ț, Diana and
Bugeag, George",
editor = "Rehm, Georg and
Bontcheva, Kalina and
Choukri, Khalid and
Haji{\v{c}}, Jan and
Piperidis, Stelios and
Vasi{\c{l}}jevs, Andrejs",
booktitle = "Proceedings of the 1st International Workshop on Language Technology Platforms",
month = may,
year = "2020",
address = "Marseille, France",
publisher = "European Language Resources Association",
url = "https://aclanthology.org/2020.iwltp-1.4/",
pages = "22--27",
language = "eng",
ISBN = "979-10-95546-64-1",
abstract = "This paper describes the on-going work carried out within the CoBiLiRo (Bimodal Corpus for Romanian Language) research project, part of ReTeRom (Resources and Technologies for Developing Human-Machine Interfaces in Romanian). Data annotation finds increasing use in speech recognition and synthesis with the goal to support learning processes. In this context, a variety of different annotation systems for application to Speech and Text Processing environments have been presented. Even if many designs for the data annotations workflow have emerged, the process of handling metadata, to manage complex user-defined annotations, is not covered enough. We propose a design of the format aimed to serve as an annotation standard for bimodal resources, which facilitates searching, editing and statistical analysis operations over it. The design and implementation of an infrastructure that houses the resources are also presented. The goal is widening the dissemination of bimodal corpora for research valorisation and use in applications. Also, this study reports on the main operations of the web Platform which hosts the corpus and the automatic conversion flows that brings the submitted files at the format accepted by the Platform."
}
<?xml version="1.0" encoding="UTF-8"?>
<modsCollection xmlns="http://www.loc.gov/mods/v3">
<mods ID="cristea-etal-2020-cobiliro">
<titleInfo>
<title>CoBiLiRo: A Research Platform for Bimodal Corpora</title>
</titleInfo>
<name type="personal">
<namePart type="given">Dan</namePart>
<namePart type="family">Cristea</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Ionuț</namePart>
<namePart type="family">Pistol</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Șerban</namePart>
<namePart type="family">Boghiu</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Anca-Diana</namePart>
<namePart type="family">Bibiri</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Daniela</namePart>
<namePart type="family">Gîfu</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Andrei</namePart>
<namePart type="family">Scutelnicu</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Mihaela</namePart>
<namePart type="family">Onofrei</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Diana</namePart>
<namePart type="family">Trandabăț</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">George</namePart>
<namePart type="family">Bugeag</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<originInfo>
<dateIssued>2020-05</dateIssued>
</originInfo>
<typeOfResource>text</typeOfResource>
<language>
<languageTerm type="text">eng</languageTerm>
</language>
<relatedItem type="host">
<titleInfo>
<title>Proceedings of the 1st International Workshop on Language Technology Platforms</title>
</titleInfo>
<name type="personal">
<namePart type="given">Georg</namePart>
<namePart type="family">Rehm</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Kalina</namePart>
<namePart type="family">Bontcheva</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Khalid</namePart>
<namePart type="family">Choukri</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Jan</namePart>
<namePart type="family">Hajič</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Stelios</namePart>
<namePart type="family">Piperidis</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Andrejs</namePart>
<namePart type="family">Vasiļjevs</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<originInfo>
<publisher>European Language Resources Association</publisher>
<place>
<placeTerm type="text">Marseille, France</placeTerm>
</place>
</originInfo>
<genre authority="marcgt">conference publication</genre>
<identifier type="isbn">979-10-95546-64-1</identifier>
</relatedItem>
<abstract>This paper describes the on-going work carried out within the CoBiLiRo (Bimodal Corpus for Romanian Language) research project, part of ReTeRom (Resources and Technologies for Developing Human-Machine Interfaces in Romanian). Data annotation finds increasing use in speech recognition and synthesis with the goal to support learning processes. In this context, a variety of different annotation systems for application to Speech and Text Processing environments have been presented. Even if many designs for the data annotations workflow have emerged, the process of handling metadata, to manage complex user-defined annotations, is not covered enough. We propose a design of the format aimed to serve as an annotation standard for bimodal resources, which facilitates searching, editing and statistical analysis operations over it. The design and implementation of an infrastructure that houses the resources are also presented. The goal is widening the dissemination of bimodal corpora for research valorisation and use in applications. Also, this study reports on the main operations of the web Platform which hosts the corpus and the automatic conversion flows that brings the submitted files at the format accepted by the Platform.</abstract>
<identifier type="citekey">cristea-etal-2020-cobiliro</identifier>
<location>
<url>https://aclanthology.org/2020.iwltp-1.4/</url>
</location>
<part>
<date>2020-05</date>
<extent unit="page">
<start>22</start>
<end>27</end>
</extent>
</part>
</mods>
</modsCollection>
%0 Conference Proceedings
%T CoBiLiRo: A Research Platform for Bimodal Corpora
%A Cristea, Dan
%A Pistol, Ionuț
%A Boghiu, Șerban
%A Bibiri, Anca-Diana
%A Gîfu, Daniela
%A Scutelnicu, Andrei
%A Onofrei, Mihaela
%A Trandabăț, Diana
%A Bugeag, George
%Y Rehm, Georg
%Y Bontcheva, Kalina
%Y Choukri, Khalid
%Y Hajič, Jan
%Y Piperidis, Stelios
%Y Vasiļjevs, Andrejs
%S Proceedings of the 1st International Workshop on Language Technology Platforms
%D 2020
%8 May
%I European Language Resources Association
%C Marseille, France
%@ 979-10-95546-64-1
%G eng
%F cristea-etal-2020-cobiliro
%X This paper describes the on-going work carried out within the CoBiLiRo (Bimodal Corpus for Romanian Language) research project, part of ReTeRom (Resources and Technologies for Developing Human-Machine Interfaces in Romanian). Data annotation finds increasing use in speech recognition and synthesis with the goal to support learning processes. In this context, a variety of different annotation systems for application to Speech and Text Processing environments have been presented. Even if many designs for the data annotations workflow have emerged, the process of handling metadata, to manage complex user-defined annotations, is not covered enough. We propose a design of the format aimed to serve as an annotation standard for bimodal resources, which facilitates searching, editing and statistical analysis operations over it. The design and implementation of an infrastructure that houses the resources are also presented. The goal is widening the dissemination of bimodal corpora for research valorisation and use in applications. Also, this study reports on the main operations of the web Platform which hosts the corpus and the automatic conversion flows that brings the submitted files at the format accepted by the Platform.
%U https://aclanthology.org/2020.iwltp-1.4/
%P 22-27
Markdown (Informal)
[CoBiLiRo: A Research Platform for Bimodal Corpora](https://aclanthology.org/2020.iwltp-1.4/) (Cristea et al., IWLTP 2020)
ACL
- Dan Cristea, Ionuț Pistol, Șerban Boghiu, Anca-Diana Bibiri, Daniela Gîfu, Andrei Scutelnicu, Mihaela Onofrei, Diana Trandabăț, and George Bugeag. 2020. CoBiLiRo: A Research Platform for Bimodal Corpora. In Proceedings of the 1st International Workshop on Language Technology Platforms, pages 22–27, Marseille, France. European Language Resources Association.