[{"@type":"PropertyValue","name":"Format","value":"16kHz, 16 bit, wav, mono channel"},{"@type":"PropertyValue","name":"Age","value":"12 years old and younger children"},{"@type":"PropertyValue","name":"Recording environment","value":"Low background noise"},{"@type":"PropertyValue","name":"Country","value":"South Korea(KOR)"},{"@type":"PropertyValue","name":"Language(Region) Code","value":"ko-KR"},{"@type":"PropertyValue","name":"Language","value":"Korean"},{"@type":"PropertyValue","name":"Features of annotation","value":"Transcription text, timestamp, speaker ID, gender, noise"},{"@type":"PropertyValue","name":"Accuracy","value":"Word Accuracy Rate (WAR) 98%"}]
{"id":1329,"datatype":"1","titleimg":"https://www.nexdata.ai/shujutang/static/image/index/datatang_yuyin_default.webp","type1":"165","type1str":null,"type2":"166","type2str":null,"dataname":"93 Hours Korean Children Speech Dataset for Speech Recognition","datazy":[{"title":"Format","content":"16kHz, 16 bit, wav, mono channel","desc":"Format"},{"title":"Age","content":"12 years old and younger children","desc":"Age"},{"title":"Recording environment","content":"Low background noise","desc":"Recording environment"},{"title":"Country","content":"South Korea(KOR)","desc":"Country"},{"title":"Language(Region) Code","content":"ko-KR","desc":"Language(Region) Code"},{"title":"Language","content":"Korean","desc":"Language"},{"title":"Features of annotation","content":"Transcription text, timestamp, speaker ID, gender, noise","desc":"Features of annotation"},{"title":"Accuracy","content":"Word Accuracy Rate (WAR) 98%","desc":"Accuracy"}],"datatag":"Korean,Casual Conversation,Monologue,Asr","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_9.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_9.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=nEDlDYd3%2BCxORj%2B5rm2IwG9kmJI%3D","intro":"우와, 병아리다. [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_5.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_5.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=U%2F2aAqGFwUZjMEyJeA7AouVbe%2Bs%3D","intro":"안돼! 내가 줘! [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_3.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_3.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=CwO48%2B6MlKQdRAzCf91zXYL346g%3D","intro":"나 엄마랑 노는 거 재미있지만 [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_8.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_8.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=3Uusq0kD1XUZu%2BkOx%2FJ7N4jpsEg%3D","intro":"이게 뭐예요? [N]","size":0,"progress":100,"type":"mp3"},{"name":"/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_1.wav","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY231031006_demo1730455202662/APY231031006_demo/KR000029_1.wav?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=vcALPhzQvrPlGbKxNcs3ScG8i6M%3D","intro":"[OVERLAP/]안녕? 난 나영이야.[/OVERLAP][N]","size":0,"progress":100,"type":"mp3"}],"officialSummary":"This 93 hours Korean children speech dataset mirrors real-world interactions. Each audio sample is transcribed and includes text content, speaker ID, gender, age, accent, and other attributes.Our dataset, collected from a wide and diverse range of speakers (children aged 12 and under), improves the model's performance on real and complex tasks by geographic location. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.","dataexampl":null,"datakeyword":["Korean children speech dataset","Korean child speech dataset","Korean children ASR dataset","Korean children conversational speech"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Data Type,Language","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"speechRec","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"93 Hours Korean(Korea) Children Real-world Casual Conversation and Monologue speech dataset","BGimg":"brightSpot_audio","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
93 Hours Korean Children Speech Dataset for Speech Recognition
Korean children speech dataset
Korean child speech dataset
Korean children ASR dataset
Korean children conversational speech
This 93 hours Korean children speech dataset mirrors real-world interactions. Each audio sample is transcribed and includes text content, speaker ID, gender, age, accent, and other attributes.Our dataset, collected from a wide and diverse range of speakers (children aged 12 and under), improves the model's performance on real and complex tasks by geographic location. Quality tested by various AI companies. We strictly adhere to data protection regulations and privacy standards, ensuring the maintenance of user privacy and legal rights throughout the data collection, storage, and usage processes, our datasets are all GDPR, CCPA, PIPL complied.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.