[{"@type":"PropertyValue","name":"Storage format","value":"TXT"},{"@type":"PropertyValue","name":"Data content","value":"Chinese-Polish Parallel Corpus Data, content has been preliminarily categorized, covering the fields of technology, healthcare, tourism, spoken, news and military."},{"@type":"PropertyValue","name":"Data size","value":"1.99 million pairs of Chinese-Polish Parallel Corpus Data."},{"@type":"PropertyValue","name":"Language","value":"Chinese, Polish"},{"@type":"PropertyValue","name":"Application scenario","value":"machine translation"}]
{"id":1337,"datatype":"1","titleimg":"https://www.nexdata.ai/shujutang/static/image/index/datatang_wenben_default.webp","type1":"183","type1str":null,"type2":"185","type2str":null,"dataname":"1.98M Chinese-Polish Sentence Pairs – Machine Translation Dataset","datazy":[{"title":"Storage format","content":"TXT","desc":"Storage format"},{"title":"Data content","content":"Chinese-Polish Parallel Corpus Data, content has been preliminarily categorized, covering the fields of technology, healthcare, tourism, spoken, news and military.","desc":"Data content"},{"title":"Data size","content":"1.99 million pairs of Chinese-Polish Parallel Corpus Data.","desc":"Data size"},{"title":"Language","content":"Chinese, Polish","desc":"Language"},{"title":"Application scenario","content":"machine translation","desc":"Application scenario"}],"datatag":"Chinese,Polish,Parallel","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/all_demo%2FAPY230901002_demo1729159203287/zh_pl_demo.png","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/all_demo%252FAPY230901002_demo1729159203287/zh_pl_demo.png?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=uL9efmbp4GOqp9AA%2F6opWFsJmmk%3D","intro":"","size":0,"progress":100,"type":"jpg"}],"officialSummary":"This dataset contains 1.98 million Chinese-Polish parallel sentence pairs stored in TXT format. The data has undergone cleaning, anonymization, and quality inspection to improve data quality and protect sensitive information. The dataset is suitable for machine translation, bilingual language modeling, and translation model training.","dataexampl":null,"datakeyword":["Chinese Polish parallel corpus","Chinese Polish translation dataset","Chinese Polish parallel dataset","Chinese Polish sentence pairs","Chinese Polish translation corpus"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Type","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"nlu","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,JP,PT,DE,KO,FR,ES\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"1,980,000 Groups - Chinese-Polish Parallel Corpus Data","BGimg":"","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
https://www.nexdata.ai/shujutang/static/image/index/datatang_wenben_default.webp
[{"@type":"ImageObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/all_demo%252FAPY230901002_demo1729159203287/zh_pl_demo.png?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=uL9efmbp4GOqp9AA%2F6opWFsJmmk%3D"}]
1.98M Chinese-Polish Sentence Pairs – Machine Translation Dataset
Chinese Polish parallel corpus
Chinese Polish translation dataset
Chinese Polish parallel dataset
Chinese Polish sentence pairs
Chinese Polish translation corpus
This dataset contains 1.98 million Chinese-Polish parallel sentence pairs stored in TXT format. The data has undergone cleaning, anonymization, and quality inspection to improve data quality and protect sensitive information. The dataset is suitable for machine translation, bilingual language modeling, and translation model training.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.
![Specifications]()
Specifications
Data content
Chinese-Polish Parallel Corpus Data, content has been preliminarily categorized, covering the fields of technology, healthcare, tourism, spoken, news and military.
Data size
1.99 million pairs of Chinese-Polish Parallel Corpus Data.
Application scenario
machine translation
![Sample]()
Sample
![Recommended Datasets]()
Recommended Dataset
Tell Us Your Special Needs
61a9850d-bf18-4e15-8067-b49326df92d4