[{"@type":"PropertyValue","name":"Storage format","value":"TXT"},{"@type":"PropertyValue","name":"Data content","value":"Chinese-Tibetan Parallel Corpus Data"},{"@type":"PropertyValue","name":"Data size","value":"5.01 million pairs of Chinese-Tibetan Parallel Corpus Data. The Chinese sentences contain 20.8 characters on average."},{"@type":"PropertyValue","name":"Language","value":"Chinese, Tibetan"},{"@type":"PropertyValue","name":"Application scenario","value":"machine translation"}]
{"id":1236,"datatype":"1","titleimg":"https://www.nexdata.ai/shujutang/static/image/index/datatang_wenben_default.webp","type1":"183","type1str":null,"type2":"185","type2str":null,"dataname":"5.01M Chinese-Tibetan Sentence Pairs – Machine Translation Dataset","datazy":[{"title":"Storage format","desc":"Storage format","content":"TXT"},{"title":"Data content","desc":"Data content","content":"Chinese-Tibetan Parallel Corpus Data"},{"title":"Data size","desc":"Data size","content":"5.01 million pairs of Chinese-Tibetan Parallel Corpus Data. The Chinese sentences contain 20.8 characters on average."},{"title":"Language","desc":"Language","content":"Chinese, Tibetan"},{"title":"Application scenario","desc":"Application scenario","content":"machine translation"}],"datatag":"Chinese,Tibetan,Chinese-Tibetan,Parallel Corpus","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY230315001_demo1729159200808/demo.png","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230315001_demo1729159200808/demo.png?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=tLL0sffQZBePZWEDJDuUT0Q%2B7oI%3D","intro":"","size":0,"progress":100,"type":"jpg"}],"officialSummary":"This dataset contains 5.01 million Chinese-Tibetan parallel sentence pairs stored in TXT format. The data has undergone cleaning, anonymization, and quality inspection, which can be used as a basic corpus for text data analysis and in fields such as machine translation.","dataexampl":null,"datakeyword":["Chinese Tibetan parallel corpus","Chinese Tibetan translation dataset","Chinese Tibetan parallel dataset","Chinese Tibetan sentence pairs","Chinese Tibetan translation corpus"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Type","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"nlu","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"5,010,000 Groups - Chinese-Tibetan Parallel Corpus Data","BGimg":"","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
https://www.nexdata.ai/shujutang/static/image/index/datatang_wenben_default.webp
[{"@type":"ImageObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY230315001_demo1729159200808/demo.png?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=tLL0sffQZBePZWEDJDuUT0Q%2B7oI%3D"}]
5.01M Chinese-Tibetan Sentence Pairs – Machine Translation Dataset
Chinese Tibetan parallel corpus
Chinese Tibetan translation dataset
Chinese Tibetan parallel dataset
Chinese Tibetan sentence pairs
Chinese Tibetan translation corpus
This dataset contains 5.01 million Chinese-Tibetan parallel sentence pairs stored in TXT format. The data has undergone cleaning, anonymization, and quality inspection, which can be used as a basic corpus for text data analysis and in fields such as machine translation.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.
![Specifications]()
Specifications
Data content
Chinese-Tibetan Parallel Corpus Data
Data size
5.01 million pairs of Chinese-Tibetan Parallel Corpus Data. The Chinese sentences contain 20.8 characters on average.
Application scenario
machine translation
![Sample]()
Sample
![Recommended Datasets]()
Recommended Dataset
Tell Us Your Special Needs
5fcd9e05-947c-4591-9caa-60b49deb3c1d