[{"@type":"PropertyValue","name":"Format","value":"48,000Hz, 24bit, uncompressed wav, mono channel;"},{"@type":"PropertyValue","name":"Recording environment","value":"professional recording studio;"},{"@type":"PropertyValue","name":"Recording content","value":"customer service;"},{"@type":"PropertyValue","name":"Speaker","value":"Brazilian;"},{"@type":"PropertyValue","name":"Annotation","value":"word and phoneme transcription, prosodic boundary annotation;"},{"@type":"PropertyValue","name":"Device","value":"microphone;"},{"@type":"PropertyValue","name":"Language","value":"Brazilian Portuguese;"},{"@type":"PropertyValue","name":"Application scenarios","value":"speech synthesis."}]
{"id":1895,"datatype":"1","titleimg":"https://www.nexdata.ai/shujutang/static/image/index/datatang_yuyin_default.webp","type1":"165","type1str":null,"type2":"219","type2str":null,"dataname":"10 Hour Brazilian Portuguese Speech Synthesis Dataset for TTS Training","datazy":[{"title":"Format","content":"48,000Hz, 24bit, uncompressed wav, mono channel;"},{"title":"Recording environment","content":"professional recording studio;"},{"title":"Recording content","content":"customer service;"},{"title":"Speaker","content":"Brazilian;"},{"title":"Annotation","content":"word and phoneme transcription, prosodic boundary annotation;"},{"title":"Device","content":"microphone;"},{"title":"Language","content":"Brazilian Portuguese;"},{"title":"Application scenarios","content":"speech synthesis."}],"datatag":"Brazilian Portuguese,TTS,Female","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":null,"samplePresentation":[{"name":"000261.wav","url":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000261.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=63wsnEV%2F3EkCf9%2F45q2khpAzyc0%3D","intro":"S I1 / CH I0 . V EH1 R / D U1 . V I . D AX0 S / S O1 . B R I0 / U S / SH AXN0 . P U1 S / K ON0 . JH I0 . S I0 . O0 . N AX0 . D O1 . R I S / O W0 / T R AX0 . T AX0 . M EN1 . T U S / K AX0 . P I0 . L A1 . R I S / N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A W0 / K L I0 . EN1 . CH I0 / I S . T A1 / JH I S . P O0 . N I1 . V E W / P A1 . R AX0 / F O R . N EH0 . S EH1 R / IN0 . F O R . M AX0 . S ON1 JN S / S O1 . B R I0 / IN0 . G R EH0 . JH I0 . EN1 . CH I S / U1","size":1774842,"progress":100,"type":"mp3"},{"name":"000445.wav","url":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000445.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=55W78CiEj1%2BGOhbj6WrTZpBSNSw%3D","intro":"S I / P R EH . S I . Z A1 R / JH I / IN . F O R . M AX . S ON1 JN S / S O1 . B R I / A / S EH . L EH . S AXN1 WN / JH I / T EN1 . D AX S / K ON / I S . P A1 . S U / P A1 . R AX / M U1 W . CH I . P L U S / P EH1 CH I S / I . K I . P AX . M EN1 . T U S / P A1 . R AX / AX . N I . M A1 J S / I / O R . G AX . N I . Z AX . S AXN1 WN / JH I / AX . CH I . V I . D A1 . JH I S / S EH . G U1 . R AX S / P A1 . R AX / F AX . M I1 . L I . A S / K ON / M U1 W . CH I . P L U S / P EH1 CH I S / N AO1 . S AX / I","size":2182212,"progress":100,"type":"mp3"},{"name":"000789.wav","url":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000789.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=xrnCk7mYP7tZlZ2eIOfGhR%2B7XGE%3D","intro":"N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A1 W / K L I0 . EN1 . CH I0 / JH I0 / L AO1 . ZH AX0 / JH I0 / D EH0 . K O0 . R AX0 . S AXN1 WN / P A1 . R AX0 / K A1 . Z AX0 / I0 S . T A1 / JH I0 S . P O0 . N I1 . V E W / P A1 . R AX0 / X EH0 S . P ON0 . D EH1 R / A0 / P EH R . G UN1 . T AX0 S / S O1 . B R I0 / A0 / IN0 S . T AX0 . L AX0 . S AXN1 WN / JH I0 / K W A1 . D R U0 S / I0 / I0 S . P EH1 . LJ U0 S","size":1397268,"progress":100,"type":"mp3"},{"name":"000841.wav","url":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000841.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=cPtwrc8yM4xK5nnSzW1AVgek4Dg%3D","intro":"S I1 / V O0 . S EH1 / D EH0 . Z EH1 . ZH AX0 / T EH1 R / UN0 / P EH0 R . S O0 . N AX0 W / T R A1 J . N E R / N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A W1 / K L I0 . EN1 . CH I0 / JH I0 / AX0 . K AX0 . D EH0 . M I1 . AX0 / AX0 . P R EH0 . Z EN0 . T AX0 . R A1 / U0 S / P R O0 . F I0 . S I0 . O0 . N A1 J S / JH I0 S . P O0 . N I1 . V E I S / I0 / S U1 . A0 Z / I0 S . P EH0 . S I0 . A0 . L I0 . D A1 . JH I S","size":1384092,"progress":100,"type":"mp3"},{"name":"000851.wav","url":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000851.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=wC9df6pWtLsj9XKevPP1I0QjAD4%3D","intro":"S I1 / V O0 . S EH1 / D EH0 . Z EH1 . ZH AX0 / P R O0 . T EH0 . ZH EH1 R / A0 / S AX0 . U1 . JH I0 / D U0 / S EH1 U0 / P EH1 T 0 / N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A1 W / K L I0 . EN1 . CH I0 / JH I0 / P L A1 . N U0 S / JH I0 / S AX0 . U1 . JH I0 / AX0 . N I0 . M A1 W / AX0 . P R EH0 . Z EN0 . T AX0 . R A1 / A1 Z / O0 P . S ON1 JN S / JH I0 / K O0 . B EH R . T U1 . R AX0","size":1367744,"progress":100,"type":"mp3"}],"officialSummary":"This dataset contains 10 hours of Brazilian Portuguese speech recordings, collected from native Brazilian speakers. The corpus is related to the customer service field. The dataset features balanced phoneme coverage. Professional phonetician participates in the annotation. It precisely matches with the research and development needs of the speech synthesis.","dataexampl":null,"datakeyword":["Portuguese TTS dataset","customer service speech dataset","call center speech corpus","TTS training dataset Brazilian Portuguese","Brazilian Portuguese speech synthesis dataset"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Language,Voice Type","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechSyn","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN\"},{\"code\":\"3\",\"language\":\"EN\"}]","productNameEn":"10 Hours - Brazilian Portuguese Speech Synthesis Corpus","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
https://www.nexdata.ai/shujutang/static/image/index/datatang_yuyin_default.webp
[{"@type":"AudioObject","embedUrl":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000261.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=63wsnEV%2F3EkCf9%2F45q2khpAzyc0%3D"},{"@type":"AudioObject","embedUrl":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000445.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=55W78CiEj1%2BGOhbj6WrTZpBSNSw%3D"},{"@type":"AudioObject","embedUrl":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000789.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=xrnCk7mYP7tZlZ2eIOfGhR%2B7XGE%3D"},{"@type":"AudioObject","embedUrl":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000841.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=cPtwrc8yM4xK5nnSzW1AVgek4Dg%3D"},{"@type":"AudioObject","embedUrl":"https://storage-product.datatang.com/damp/product/instructions_zh/20260119161907/000851.wav?Expires=4102415999&OSSAccessKeyId=LTAI5tEBeSWUJiqjXvBMsxEu&Signature=wC9df6pWtLsj9XKevPP1I0QjAD4%3D"}]
10 Hour Brazilian Portuguese Speech Synthesis Dataset for TTS Training
Portuguese TTS dataset
customer service speech dataset
call center speech corpus
TTS training dataset Brazilian Portuguese
Brazilian Portuguese speech synthesis dataset
This dataset contains 10 hours of Brazilian Portuguese speech recordings, collected from native Brazilian speakers. The corpus is related to the customer service field. The dataset features balanced phoneme coverage. Professional phonetician participates in the annotation. It precisely matches with the research and development needs of the speech synthesis.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.
![Specifications]()
Specifications
Format
48,000Hz, 24bit, uncompressed wav, mono channel;
Recording environment
professional recording studio;
Recording content
customer service;
Annotation
word and phoneme transcription, prosodic boundary annotation;
Language
Brazilian Portuguese;
Application scenarios
speech synthesis.
![Sample]()
Sample
Audio
S I1 / CH I0 . V EH1 R / D U1 . V I . D AX0 S / S O1 . B R I0 / U S / SH AXN0 . P U1 S / K ON0 . JH I0 . S I0 . O0 . N AX0 . D O1 . R I S / O W0 / T R AX0 . T AX0 . M EN1 . T U S / K AX0 . P I0 . L A1 . R I S / N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A W0 / K L I0 . EN1 . CH I0 / I S . T A1 / JH I S . P O0 . N I1 . V E W / P A1 . R AX0 / F O R . N EH0 . S EH1 R / IN0 . F O R . M AX0 . S ON1 JN S / S O1 . B R I0 / IN0 . G R EH0 . JH I0 . EN1 . CH I S / U1
Audio
S I / P R EH . S I . Z A1 R / JH I / IN . F O R . M AX . S ON1 JN S / S O1 . B R I / A / S EH . L EH . S AXN1 WN / JH I / T EN1 . D AX S / K ON / I S . P A1 . S U / P A1 . R AX / M U1 W . CH I . P L U S / P EH1 CH I S / I . K I . P AX . M EN1 . T U S / P A1 . R AX / AX . N I . M A1 J S / I / O R . G AX . N I . Z AX . S AXN1 WN / JH I / AX . CH I . V I . D A1 . JH I S / S EH . G U1 . R AX S / P A1 . R AX / F AX . M I1 . L I . A S / K ON / M U1 W . CH I . P L U S / P EH1 CH I S / N AO1 . S AX / I
Audio
N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A1 W / K L I0 . EN1 . CH I0 / JH I0 / L AO1 . ZH AX0 / JH I0 / D EH0 . K O0 . R AX0 . S AXN1 WN / P A1 . R AX0 / K A1 . Z AX0 / I0 S . T A1 / JH I0 S . P O0 . N I1 . V E W / P A1 . R AX0 / X EH0 S . P ON0 . D EH1 R / A0 / P EH R . G UN1 . T AX0 S / S O1 . B R I0 / A0 / IN0 S . T AX0 . L AX0 . S AXN1 WN / JH I0 / K W A1 . D R U0 S / I0 / I0 S . P EH1 . LJ U0 S
Audio
S I1 / V O0 . S EH1 / D EH0 . Z EH1 . ZH AX0 / T EH1 R / UN0 / P EH0 R . S O0 . N AX0 W / T R A1 J . N E R / N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A W1 / K L I0 . EN1 . CH I0 / JH I0 / AX0 . K AX0 . D EH0 . M I1 . AX0 / AX0 . P R EH0 . Z EN0 . T AX0 . R A1 / U0 S / P R O0 . F I0 . S I0 . O0 . N A1 J S / JH I0 S . P O0 . N I1 . V E I S / I0 / S U1 . A0 Z / I0 S . P EH0 . S I0 . A0 . L I0 . D A1 . JH I S
Audio
S I1 / V O0 . S EH1 / D EH0 . Z EH1 . ZH AX0 / P R O0 . T EH0 . ZH EH1 R / A0 / S AX0 . U1 . JH I0 / D U0 / S EH1 U0 / P EH1 T 0 / N AO1 . S AX0 / I0 . K I1 . P I0 / JH I0 / AX0 . T EN0 . JH I0 . M EN1 . T U0 / A1 W / K L I0 . EN1 . CH I0 / JH I0 / P L A1 . N U0 S / JH I0 / S AX0 . U1 . JH I0 / AX0 . N I0 . M A1 W / AX0 . P R EH0 . Z EN0 . T AX0 . R A1 / A1 Z / O0 P . S ON1 JN S / JH I0 / K O0 . B EH R . T U1 . R AX0
![Recommended Datasets]()
Recommended Dataset
Tell Us Your Special Needs

What languages and voice characteristics are covered by Nexdata’s speech synthesis datasets?

Nexdata offers speech synthesis datasets covering a broad range of languages, dialects, accents, and voice types, supported by extensive global language resources. Our datasets include diverse speakers, speaking styles, emotions, and recording scenarios to support natural and expressive Text-to-Speech (TTS) model development.

Can Nexdata customize speech synthesis datasets for specific languages or requirements?

Yes. If our off-the-shelf TTS datasets do not fully meet your requirements, Nexdata provides flexible custom data collection, transcription, annotation, and quality control services. We can customize datasets based on target languages or dialects, speaker profiles, voice characteristics, emotions, speaking styles, recording environments, and data volume.

How does Nexdata ensure the quality and scalability of speech synthesis datasets?

Nexdata applies multi-stage quality control throughout voice data collection, transcription, annotation, and validation. Combined with our extensive language resources and scalable collection capabilities, we can support both large-scale multilingual TTS projects and specialized datasets for specific voices, accents, emotions, or speech scenarios.
7fdabfcd-6efc-4079-ae33-1fe2405a856b