[{"@type":"PropertyValue","name":"Format","value":"16kHz, 16bit, uncompressed wav, mono channel;"},{"@type":"PropertyValue","name":"Recording condition","value":"Low background noise (indoor);"},{"@type":"PropertyValue","name":"Content category","value":"economy, entertainment, news, informal language, numbers, alphabet;"},{"@type":"PropertyValue","name":"Recording device","value":"Android Smartphone:iPhone = 3.5:1;"},{"@type":"PropertyValue","name":"Speaker","value":"352 people in total, from Spain, Mexico, Venezuela and other countries, 55% male and 45% female;"},{"@type":"PropertyValue","name":"Country","value":"Spain(ESP), Mexico(MEX), Venezuela(VEN), etc.;"},{"@type":"PropertyValue","name":"Language","value":"Spanish;"},{"@type":"PropertyValue","name":"Features of annotation","value":"Transcription text; timestamp; 5 noise symbols; special identifiers"},{"@type":"PropertyValue","name":"Accuracy Rate","value":"Sentence Accuracy Rate (SAR) 95%"}]
{"id":116,"datatype":"1","titleimg":"https://res.datatang.com/asset/productNew/APY161101034_R.png?Expires=2007353626&OSSAccessKeyId=LTAI5tQwXnJZbubgVfVa1ep9&Signature=3n%2BkydGWNBHjs8lurT/ACT8ErJo%3D","type1":"165","type1str":null,"type2":"166","type2str":null,"dataname":"227 Hours Multi-Accent Spanish Speech Dataset with Transcripts for ASR Training","datazy":[{"title":"Format","desc":"Format","content":"16kHz, 16bit, uncompressed wav, mono channel;"},{"title":"Recording condition","desc":"Recording condition","content":"Low background noise (indoor);"},{"title":"Content category","desc":"Content category","content":"economy, entertainment, news, informal language, numbers, alphabet;"},{"title":"Recording device","desc":"Recording device","content":"Android Smartphone:iPhone = 3.5:1;"},{"title":"Speaker","desc":"Speaker","content":"352 people in total, from Spain, Mexico, Venezuela and other countries, 55% male and 45% female;"},{"title":"Country","desc":"Country","content":"Spain(ESP), Mexico(MEX), Venezuela(VEN), etc.;"},{"title":"Language","desc":"Language","content":"Spanish;"},{"title":"Features of annotation","desc":"Features of annotation","content":"Transcription text; timestamp; 5 noise symbols; special identifiers"},{"title":"Accuracy Rate","desc":"Accuracy Rate","content":"Sentence Accuracy Rate (SAR) 95%"}],"datatag":"Spanish,Smartphone,Reading,Scripted Monologue","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":["227 hours 352 speakers","sentence accuracy rate 95%"],"samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/dataDemo%2FAPY161101034_R1695808872472/T0108G0318S0018.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/dataDemo%252FAPY161101034_R1695808872472/T0108G0318S0018.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=7GpTkwXAX%2BbkxR%2BvxpRDvpuvCwU%3D","intro":"Esta revolucionaria experiencia de compra se encuentra entre las primeras soluciones de este tipo para la televisión","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/dataDemo%2FAPY161101034_R1695808872472/T9108G0248Q0051.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/dataDemo%252FAPY161101034_R1695808872472/T9108G0248Q0051.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=ekhFSMgD8RQJ7%2F8luwevWMrv%2F%2B8%3D","intro":"mil despidos en sectores relacionados con la construcción en Castilla La Mancha y Castellón Comunidad valenciana","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/dataDemo%2FAPY161101034_R1695808872472/T9108G0125Q0069.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/dataDemo%252FAPY161101034_R1695808872472/T9108G0125Q0069.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=bVNksjKKw9IEH2pXLoynVZIZGP4%3D","intro":"Esto cada día se parece más a la Cuba [/pre/] Fidel.","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/dataDemo%2FAPY161101034_R1695808872472/T9108G0113Q0055.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/dataDemo%252FAPY161101034_R1695808872472/T9108G0113Q0055.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=5qE6exgOSzLE%2Bk7sgna%2B1yjUeQU%3D","intro":"a los nacionales cualificados nos los quedamos,que este es su país","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/dataDemo%2FAPY161101034_R1695808872472/T9108G0125Q0063.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/dataDemo%252FAPY161101034_R1695808872472/T9108G0125Q0063.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=ddp52bTw9X2cA6lQKf61SeO7RKs%3D","intro":"El despido libre sería lo mejor no % no soy empresario,soy trabajador","size":0,"progress":100,"type":"mp3"}],"officialSummary":"This dataset contains 227 hours of Spanish scripted speech recorded as monologues based on predefined texts. It includes recordings from 352 native speakers from Spain, Mexico, Venezuela, and other Spanish-speaking regions. The speech content covers economics, entertainment, news, informal language, numbers, alphabet sequences, and other general domains. Each recording includes high-quality transcripts, timestamps, noise labels, and additional speaker metadata. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.","dataexampl":null,"datakeyword":["spanish speech to text dataset","spanish asr dataset","spanish speech recognition dataset","spanish speech dataset for asr training","smartphone speech dataset"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Data Type,Language","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechRec","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"227 Hours - Spanish Speech Data by Mobile Phone_R","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
227 Hours Multi-Accent Spanish Speech Dataset with Transcripts for ASR Training
spanish speech to text dataset
spanish asr dataset
spanish speech recognition dataset
spanish speech dataset for asr training
smartphone speech dataset
This dataset contains 227 hours of Spanish scripted speech recorded as monologues based on predefined texts. It includes recordings from 352 native speakers from Spain, Mexico, Venezuela, and other Spanish-speaking regions. The speech content covers economics, entertainment, news, informal language, numbers, alphabet sequences, and other general domains. Each recording includes high-quality transcripts, timestamps, noise labels, and additional speaker metadata. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.