[{"@type":"PropertyValue","name":"Format","value":"16kHz, 16 bit, wav, mono channel"},{"@type":"PropertyValue","name":"Age","value":"12 years old and younger children"},{"@type":"PropertyValue","name":"Recording environment","value":"Low background noise"},{"@type":"PropertyValue","name":"Country","value":"Portugal(PT)"},{"@type":"PropertyValue","name":"Language(Region) Code","value":"pt-PT"},{"@type":"PropertyValue","name":"Language","value":"Portuguese"},{"@type":"PropertyValue","name":"Features of annotation","value":"Transcription text, timestamp, speaker ID, gender, noise"},{"@type":"PropertyValue","name":"Accuracy","value":"Word Accuracy Rate (WAR) 98%"}]
{"id":1325,"datatype":"1","titleimg":"https://www.nexdata.ai/shujutang/static/image/index/datatang_yuyin_default.webp","type1":"165","type1str":null,"type2":"166","type2str":null,"dataname":"58 Hours European Portuguese Children Speech Dataset for ASR Training","datazy":[{"title":"Format","content":"16kHz, 16 bit, wav, mono channel","desc":"Format"},{"title":"Age","content":"12 years old and younger children","desc":"Age"},{"title":"Recording environment","content":"Low background noise","desc":"Recording environment"},{"title":"Country","content":"Portugal(PT)","desc":"Country"},{"title":"Language(Region) Code","content":"pt-PT","desc":"Language(Region) Code"},{"title":"Language","content":"Portuguese","desc":"Language"},{"title":"Features of annotation","content":"Transcription text, timestamp, speaker ID, gender, noise","desc":"Features of annotation"},{"title":"Accuracy","content":"Word Accuracy Rate (WAR) 98%","desc":"Accuracy"}],"datatag":"Portuguese,Casual Conversation,Monologue,Asr","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000065_6.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000065_6.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=EzgwnYaC3mwrvHlN%2BGJEa%2FQuiFY%3D","intro":"Uma caneta preta para desenhar, é o que vocês querem, [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000065_2.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000065_2.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=LIEdak1WKuJF2ILJdONQvj9oios%3D","intro":"Sim, porque eu estou aqui para mais um vídeo e esta vez estou aqui bem feliz porque o vídeo de hoje é[N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000065_10.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000065_10.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=mpo06XGKd832oC2NDD7HIKv%2FMSs%3D","intro":"Já vi este vídeo super fixe! Então, vou te ver como é que se faz bem, tá?[N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000005_11.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000005_11.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=nFH%2Fm4Fm7HQMDS%2F419u5O5AWZeI%3D","intro":"[OVERLAP/] Nós até viramos [/OVERLAP] aqui para vocês verem. [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000005_4.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230830003_demo1711101638113/APY230830003_demo/000005_4.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=84jDl0FbQpOJ819uFMSGFfHQWBA%3D","intro":"[OVERLAP/] Fui eu que dei o nome. [/OVERLAP] [N]","size":0,"progress":100,"type":"mp3"}],"officialSummary":"This 58 hours European Portuguese children speech dataset mirrors real-world interactions. Each audio sample is transcribed and includes metadata such as speaker ID, gender, age, accent, and other relevant attributes. Our dataset, collected from a wide and diverse range of speakers (children aged 12 and under), improves the model's performance on real and complex tasks by geographic location. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.","dataexampl":null,"datakeyword":["European Portuguese children speech","child speech recognition","speech-to-text training data","transcribed speech dataset","child voice data","conversational speech corpus"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Data Type,Language","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechRec","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"58 Hours - European Portuguese Child's Spontaneous Speech Data-Nexdata","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
58 Hours European Portuguese Children Speech Dataset for ASR Training
European Portuguese children speech
child speech recognition
speech-to-text training data
transcribed speech dataset
child voice data
conversational speech corpus
This 58 hours European Portuguese children speech dataset mirrors real-world interactions. Each audio sample is transcribed and includes metadata such as speaker ID, gender, age, accent, and other relevant attributes. Our dataset, collected from a wide and diverse range of speakers (children aged 12 and under), improves the model's performance on real and complex tasks by geographic location. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.