<?xml version="1.0" encoding="UTF-8"?>

<?xml-stylesheet type="text/xsl" href="/static/oaitohtml.xsl"?>

<!--
<?xml-stylesheet type="text/xsl" href="/oaitohtml.xsl"?>
-->

<OAI-PMH xmlns="http://www.openarchives.org/OAI/2.0/" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/ http://www.openarchives.org/OAI/2.0/OAI-PMH.xsd" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance">
    <responseDate>2026-10-10T19:39:54Z</responseDate>
    <request verb="GetRecord" metadataPrefix="oai_dc" identifier="10.11922/sciencedb.j00001.00345" >https://www.scidb.cn/oai</request>
<GetRecord>
    <record>
    <header >
    <identifier>10.11922/sciencedb.j00001.00345</identifier>
    <datestamp>2022-05-07T14:34:36Z</datestamp>
</header>
    <metadata>
        
<oai_dc:dc xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:dc="http://purl.org/dc/elements/1.1/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
  <dc:date>2022-05-07</dc:date>
  <dc:title>A dataset of Mongolian-Chinese speech translation</dc:title>
  <dc:identifier>doi:10.11922/sciencedb.j00001.00345</dc:identifier>
  <dc:language>en</dc:language>
  <dc:description>Due to the lack of public datasets, few researches focus on speech translation in minority languages. To this end, this paper constructs a dataset of Mongolian-Chinese speech translation, named as NMLR-Mon2Chs ST. The dataset consists of Mongolian speech, Mongolian and Chinese text. First, Mongolian speech were obtained from 36 Mongols aged between 20 and 25 by recording on their mobile phones. Then, the corresponding Chinese texts were annotated by professionals. In order to make sure the quality of the dataset, the preprocessing was done, such as removing the quiet speech, resampling, and normalization. As a result, a total of 25 hours of high-quality data are obtained, and the average duration of audio in the dataset is 4.2 seconds. The establishment of this dataset allows researchers access to speech translation for minority languages.</dc:description>
  <dc:subject>speech translation; Mongolian-Chinese; minority languages; low resource; dataset</dc:subject>
  <dc:creator>qi xiao ke</dc:creator>
  <dc:creator>Borjigin B.Teniger</dc:creator>
  <dc:creator>Yuan Sun</dc:creator>
  <dc:creator>Xiaobing Zhao</dc:creator>
  <dc:rights>PUBLIC</dc:rights>
  <dc:rights>https://creativecommons.org/licenses/by/4.0/</dc:rights>
  <dc:type>dataset</dc:type>
  <dc:relation>http://www.doi.org/10.11922/11-6035.csd.2021.0093.zh</dc:relation>
  <dc:publisher>Science Data Bank</dc:publisher>
</oai_dc:dc>

    </metadata>
</record>
</GetRecord>
</OAI-PMH>