[{"@type":"PropertyValue","name":"Format","value":"44,100Hz, 16bit, uncompressed wav, mono channel;"},{"@type":"PropertyValue","name":"Recording environment","value":"quiet indoor environment, low background noise, without echo;"},{"@type":"PropertyValue","name":"Recording content","value":"news and colloquial sentences;"},{"@type":"PropertyValue","name":"Speaker","value":"9 males, 29 females;"},{"@type":"PropertyValue","name":"Device","value":"microphone;"},{"@type":"PropertyValue","name":"Language","value":"Cantonese, English;"},{"@type":"PropertyValue","name":"Annotation","value":"word and phoneme transcription, prosodic boundary annotation;"},{"@type":"PropertyValue","name":"Application scenarios","value":"speech synthesis."}]
{"id":1201,"datatype":"1","titleimg":"https://res.datatang.com/asset/productNew/APY221030001.png?Expires=2007353716&OSSAccessKeyId=LTAI5tQwXnJZbubgVfVa1ep9&Signature=l0UgJBrdkzvD1ShKlGzGoeMzuZo%3D","type1":"165","type1str":null,"type2":"219","type2str":null,"dataname":"38 Speakers – Large-Scale Cantonese Speech Dataset (Hong Kong)","datazy":[{"title":"Format","content":"44,100Hz, 16bit, uncompressed wav, mono channel;","desc":"Format"},{"title":"Recording environment","content":"quiet indoor environment, low background noise, without echo;","desc":"Recording environment"},{"title":"Recording content","content":"news and colloquial sentences;","desc":"Recording content"},{"title":"Speaker","content":"9 males, 29 females;","desc":"Speaker"},{"title":"Device","content":"microphone;","desc":"Device"},{"title":"Language","content":"Cantonese, English;","desc":"Language"},{"title":"Annotation","content":"word and phoneme transcription, prosodic boundary annotation;","desc":"Annotation"},{"title":"Application scenarios","content":"speech synthesis.","desc":"Application scenarios"}],"datatag":"Tts","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/280002.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/280002.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=vzTpPeH29tbzv1%2F%2BhyLdc2Yregg%3D","intro":"280002\t咁你#1即係#1declare咗#1你有#1两份#1收入#4\tgam2 nei5 zik1 hai6 / D IH0 . K L EH1 R / zo2 nei5 jau5 loeng5 fan6 sau1 jap6","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/040001.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/040001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=f0WlHxpJU73ANV2SmfBzcmkAZZ4%3D","intro":"040001\t我#1唔太鐘意#1畀人#1話我#1cute#4\tngo5 m4 taai3 zung1 ji3 bei2 jan4 waa6 ngo5 / K Y UW1 T","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/000001.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/000001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=nSh7JP%2FROyHkTlCVn8wZyb91Au8%3D","intro":"000001\tSend嗰#1message#1俾我啊#4\tS EH1 N D / go3 / M EH1 . S IH0 JH / bei2 ngo5 aa1","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/070001.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/070001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=0Rq7PYbiSD3XG1JICjcRkyUqo4c%3D","intro":"070001\tDepends on#2你想#1揾#1咩工#4\tD IH0 . P EH1 N D Z AA1 N nei5 soeng2 wan2 me1 gung1","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/350001.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/350001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=YeBVLANuBHdII23yooMysHJTaYE%3D","intro":"350001\t識得#1佢耐#3你#1就知#1佢#1好funny#4\tsik1 dak1 keoi5 noi6 nei5 zau6 zi1 keoi5 hou2 / F AH1 . N IH0","size":0,"progress":100,"type":"mp3"}],"officialSummary":"This dataset contains recordings from 38 native Hong Kong Cantonese speakers. Professional phonetician participates in the annotation. It precisely matches with the research and development needs of the speech synthesis.","dataexampl":null,"datakeyword":["Hong Kong accent speech dataset","Cantonese TTS dataset","native Cantonese speech recordings","Cantonese audio dataset"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Language,Voice Type","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechSyn","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"}]","productNameEn":"38 People - Hong Kong Cantonese Average Tone Speech Synthesis Corpus","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
https://www.nexdata.ai/shujutang/static/image/index/datatang_yuyin_default.webp
[{"@type":"AudioObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/280002.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=vzTpPeH29tbzv1%2F%2BhyLdc2Yregg%3D"},{"@type":"AudioObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/040001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=f0WlHxpJU73ANV2SmfBzcmkAZZ4%3D"},{"@type":"AudioObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/000001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=nSh7JP%2FROyHkTlCVn8wZyb91Au8%3D"},{"@type":"AudioObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/070001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=0Rq7PYbiSD3XG1JICjcRkyUqo4c%3D"},{"@type":"AudioObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY221030001_demo1724320802860/APY221030001_demo/350001.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=YeBVLANuBHdII23yooMysHJTaYE%3D"}]
38 Speakers – Large-Scale Cantonese Speech Dataset (Hong Kong)
Hong Kong accent speech dataset
Cantonese TTS dataset
native Cantonese speech recordings
Cantonese audio dataset
This dataset contains recordings from 38 native Hong Kong Cantonese speakers. Professional phonetician participates in the annotation. It precisely matches with the research and development needs of the speech synthesis.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.
![Specifications]()
Specifications
Format
44,100Hz, 16bit, uncompressed wav, mono channel;
Recording environment
quiet indoor environment, low background noise, without echo;
Recording content
news and colloquial sentences;
Speaker
9 males, 29 females;
Language
Cantonese, English;
Annotation
word and phoneme transcription, prosodic boundary annotation;
Application scenarios
speech synthesis.
![Sample]()
Sample
Audio
280002 咁你#1即係#1declare咗#1你有#1两份#1收入#4 gam2 nei5 zik1 hai6 / D IH0 . K L EH1 R / zo2 nei5 jau5 loeng5 fan6 sau1 jap6
Audio
040001 我#1唔太鐘意#1畀人#1話我#1cute#4 ngo5 m4 taai3 zung1 ji3 bei2 jan4 waa6 ngo5 / K Y UW1 T
Audio
000001 Send嗰#1message#1俾我啊#4 S EH1 N D / go3 / M EH1 . S IH0 JH / bei2 ngo5 aa1
Audio
070001 Depends on#2你想#1揾#1咩工#4 D IH0 . P EH1 N D Z AA1 N nei5 soeng2 wan2 me1 gung1
Audio
350001 識得#1佢耐#3你#1就知#1佢#1好funny#4 sik1 dak1 keoi5 noi6 nei5 zau6 zi1 keoi5 hou2 / F AH1 . N IH0
Tell Us Your Special Needs

What languages and voice characteristics are covered by Nexdata’s speech synthesis datasets?

Nexdata offers speech synthesis datasets covering a broad range of languages, dialects, accents, and voice types, supported by extensive global language resources. Our datasets include diverse speakers, speaking styles, emotions, and recording scenarios to support natural and expressive Text-to-Speech (TTS) model development.

Can Nexdata customize speech synthesis datasets for specific languages or requirements?

Yes. If our off-the-shelf TTS datasets do not fully meet your requirements, Nexdata provides flexible custom data collection, transcription, annotation, and quality control services. We can customize datasets based on target languages or dialects, speaker profiles, voice characteristics, emotions, speaking styles, recording environments, and data volume.

How does Nexdata ensure the quality and scalability of speech synthesis datasets?

Nexdata applies multi-stage quality control throughout voice data collection, transcription, annotation, and validation. Combined with our extensive language resources and scalable collection capabilities, we can support both large-scale multilingual TTS projects and specialized datasets for specific voices, accents, emotions, or speech scenarios.
fd861eed-e878-4468-b725-31898cefaf8a