[{"@type":"PropertyValue","name":"Format","value":"16kHz, 16 bit, wav, mono channel"},{"@type":"PropertyValue","name":"Age","value":"12 years old and younger children"},{"@type":"PropertyValue","name":"Recording environment","value":"Low background noise"},{"@type":"PropertyValue","name":"Country","value":"Italy(ITA)"},{"@type":"PropertyValue","name":"Language(Region) Code","value":"it-IT"},{"@type":"PropertyValue","name":"Language","value":"Italian"},{"@type":"PropertyValue","name":"Features of annotation","value":"Transcription text, timestamp, speaker ID, gender, noise"},{"@type":"PropertyValue","name":"Accuracy","value":"Word Accuracy Rate (WAR) 98%"}]
{"id":1300,"datatype":"1","titleimg":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/asset/productNew/nexdata/APY230809002.jpg?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=EdmacVo0vIBd28fnoirypeKxNBc%3D","type1":"165","type1str":null,"type2":"166","type2str":null,"dataname":"101 Hours Italian Children Speech Dataset for ASR Training","datazy":[{"title":"Format","content":"16kHz, 16 bit, wav, mono channel","desc":"Format"},{"title":"Age","content":"12 years old and younger children","desc":"Age"},{"title":"Recording environment","content":"Low background noise","desc":"Recording environment"},{"title":"Country","content":"Italy(ITA)","desc":"Country"},{"title":"Language(Region) Code","content":"it-IT","desc":"Language(Region) Code"},{"title":"Language","content":"Italian","desc":"Language"},{"title":"Features of annotation","content":"Transcription text, timestamp, speaker ID, gender, noise","desc":"Features of annotation"},{"title":"Accuracy","content":"Word Accuracy Rate (WAR) 98%","desc":"Accuracy"}],"datatag":"Italian,Casual Conversation,Monologue,Asr","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000002_10.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000002_10.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=inE6%2BouR25hXqr3UstdpZ86QPYM%3D","intro":"Eh! [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000002_6.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000002_6.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=jnJ6CTwbkBOKbDPtioYPpf3ly%2FM%3D","intro":"No! No! Ah! [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000002_3.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000002_3.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=cBXc5yJnhsbxjwpBfXpv0qpKjUQ%3D","intro":"Con tanti passaggi segreti. [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000004_7.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000004_7.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=glrfW5ght888kDc5YQk9cGstgQY%3D","intro":"Vabbè, comincio, Filippo non accettava rimproveri da nessuno. [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000003_2.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230809002_demo1727690400589/APY230809002_demo/000003_2.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=w0%2FrMamzCvp7b5vQvhM26bYAKuI%3D","intro":"Cooper ma parli? [N]","size":0,"progress":100,"type":"mp3"}],"officialSummary":"This 101 hours Italian children speech dataset reflects real-world interactions. Each audio sample is transcribed and includes rich metadata such as speaker ID, gender, age, accent, timestamps, and noise-related attributes. Our dataset is collected from a broad and diverse group of speakers (children aged 12 and under) with wide geographical distribution, thus improving the model's performance on real-world, complex tasks. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.","dataexampl":null,"datakeyword":["Italian children speech dataset","Italian child speech corpus","Italian kids speech dataset","Italian children ASR dataset","Italian child voice dataset","Italian speech recognition training data"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Data Type,Language","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechRec","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"101 Hours - Italian(Italy) Children Real-world Casual Conversation and Monologue speech dataset","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
101 Hours Italian Children Speech Dataset for ASR Training
Italian children speech dataset
Italian child speech corpus
Italian kids speech dataset
Italian children ASR dataset
Italian child voice dataset
Italian speech recognition training data
This 101 hours Italian children speech dataset reflects real-world interactions. Each audio sample is transcribed and includes rich metadata such as speaker ID, gender, age, accent, timestamps, and noise-related attributes. Our dataset is collected from a broad and diverse group of speakers (children aged 12 and under) with wide geographical distribution, thus improving the model's performance on real-world, complex tasks. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.
This is a paid dataset licensed for commercial use. Ready-made datasets are available for immediate integration into AI projects.