[{"@type":"PropertyValue","name":"Format","value":"16kHz, 16 bit, uncompressed wav, mono channel;"},{"@type":"PropertyValue","name":"Content category","value":"Including generic domain; human-machine interaction; smart home command and control; in-car command and control; numbers;"},{"@type":"PropertyValue","name":"Recording condition","value":"Low background noise (indoor), without echo;"},{"@type":"PropertyValue","name":"Recording device","value":"Android Smartphone, iPhone;"},{"@type":"PropertyValue","name":"Speaker","value":"532 speakers totally, with 47% male and 53% female; and 59% speakers of all are in the age group of 18-25,39% speakers of all are in the age group of 26-45, 2% speakers of all are in the age group of 46-60;"},{"@type":"PropertyValue","name":"Country","value":"Portugal(PRT);"},{"@type":"PropertyValue","name":"Language","value":"English;"},{"@type":"PropertyValue","name":"Features of annotation","value":"Transcription text;"},{"@type":"PropertyValue","name":"Accuracy Rate","value":"Sentence Accuracy Rate (SAR) 95%"}]
{"id":1023,"datatype":"1","titleimg":"https://res.datatang.com/asset/productNew/APY190704003.png?Expires=2007353668&OSSAccessKeyId=LTAI5tQwXnJZbubgVfVa1ep9&Signature=fjMwUh6moGZvObbRHZPSwPest2g%3D","type1":"165","type1str":null,"type2":"166","type2str":null,"dataname":"209 Hours Portuguese Accent English Speech Dataset for ASR Training","datazy":[{"title":"Format","desc":"Format","content":"16kHz, 16 bit, uncompressed wav, mono channel;"},{"title":"Content category","desc":"Content category","content":"Including generic domain; human-machine interaction; smart home command and control; in-car command and control; numbers;"},{"title":"Recording condition","desc":"Recording condition","content":"Low background noise (indoor), without echo;"},{"title":"Recording device","desc":"Recording device","content":"Android Smartphone, iPhone;"},{"title":"Speaker","desc":"Speaker","content":"532 speakers totally, with 47% male and 53% female; and 59% speakers of all are in the age group of 18-25,39% speakers of all are in the age group of 26-45, 2% speakers of all are in the age group of 46-60;"},{"title":"Country","desc":"Country","content":"Portugal(PRT);"},{"title":"Language","desc":"Language","content":"English;"},{"title":"Features of annotation","desc":"Features of annotation","content":"Transcription text;"},{"title":"Accuracy Rate","desc":"Accuracy Rate","content":"Sentence Accuracy Rate (SAR) 95%"}],"datatag":"English,Portugal,Mobile,Read,Scripted Monologue","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":["209 hours","532 people","male and female ratio 1:1"],"samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00810S3434.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00810S3434.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=6JLFGgFvDo4Qjm97X8khpXqjY7g%3D","intro":"Please increase the speed by two","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00002S5455.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00002S5455.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=FRFSuSfLgjAvXE%2FGE9lSPXe2eWM%3D","intro":"twenty-six to four PM","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G60652S4440.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G60652S4440.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=lFKJarst592HZCdGHRQ1VshvaPM%3D","intro":"Please stop the car recorder","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00611S2283.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00611S2283.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=GHR4Lwq8XXeqsJY55FSXrNOOMYw%3D","intro":"Play the original Let's Groove version","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00002S1028.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY190704003_demo1695808957610/G00002S1028.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=HXMEJj9rj8%2FXdEp2AxLFthJbGJg%3D","intro":"The city was latecomer to the new economy, and it was relatively remotes.","size":0,"progress":100,"type":"mp3"}],"officialSummary":"This dataset contains 209 hours of scripted English speech collected from Portuguese speakers using smartphone devices. The recordings are based on predefined scripts covering general topics, human-machine interaction scenarios, smart home voice commands, automotive voice commands, numbers, and other speech applications. Each audio sample is transcribed with corresponding text content and additional metadata. Collected from 532 Portuguese speakers, the dataset captures diverse Portuguese-accented English pronunciation patterns and speaking variations, enhancing model performance in real and complex tasks.Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.","dataexampl":null,"datakeyword":["portuguese accent english speech dataset","accent speech dataset","english ASR dataset","voice assistant training data","speech recognition training data","speech command dataset"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Data Type,Language","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechRec","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"209 Hours - Portuguese Speaking English Speech Data by Mobile Phone","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
209 Hours Portuguese Accent English Speech Dataset for ASR Training
portuguese accent english speech dataset
accent speech dataset
english ASR dataset
voice assistant training data
speech recognition training data
speech command dataset
This dataset contains 209 hours of scripted English speech collected from Portuguese speakers using smartphone devices. The recordings are based on predefined scripts covering general topics, human-machine interaction scenarios, smart home voice commands, automotive voice commands, numbers, and other speech applications. Each audio sample is transcribed with corresponding text content and additional metadata. Collected from 532 Portuguese speakers, the dataset captures diverse Portuguese-accented English pronunciation patterns and speaking variations, enhancing model performance in real and complex tasks.Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.
Specifications
Format
16kHz, 16 bit, uncompressed wav, mono channel;
Content category
Including generic domain; human-machine interaction; smart home command and control; in-car command and control; numbers;
Recording condition
Low background noise (indoor), without echo;
Recording device
Android Smartphone, iPhone;
Speaker
532 speakers totally, with 47% male and 53% female; and 59% speakers of all are in the age group of 18-25,39% speakers of all are in the age group of 26-45, 2% speakers of all are in the age group of 46-60;
Country
Portugal(PRT);
Language
English;
Features of annotation
Transcription text;
Accuracy Rate
Sentence Accuracy Rate (SAR) 95%
Sample
Audio
Please increase the speed by two
Audio
twenty-six to four PM
Audio
Please stop the car recorder
Audio
Play the original Let's Groove version
Audio
The city was latecomer to the new economy, and it was relatively remotes.