[{"@type":"PropertyValue","name":"Storage format","value":"TXT"},{"@type":"PropertyValue","name":"Data content","value":"Chinese-Vietnamese Parallel Corpus Data"},{"@type":"PropertyValue","name":"Data size","value":"7.29 million pairs of Chinese-Vietnamese Parallel Corpus Data"},{"@type":"PropertyValue","name":"Language","value":"Chinese,Vietnamese"},{"@type":"PropertyValue","name":"Application scenario","value":"machine translation"},{"@type":"PropertyValue","name":"Accuracy rate","value":"90%"}]
{"id":1170,"datatype":"1","titleimg":"https://res.datatang.com/asset/productNew/APY220317001.png?Expires=2007353707&OSSAccessKeyId=LTAI5tQwXnJZbubgVfVa1ep9&Signature=LNwwKalhncI15xFg1D/IUoxTSsA%3D","type1":"183","type1str":null,"type2":"185","type2str":null,"dataname":"7.29M Chinese-Vietnamese Sentence Pairs – Machine Translation Dataset","datazy":[{"title":"Storage format","content":"TXT","desc":"Storage format"},{"title":"Data content","content":"Chinese-Vietnamese Parallel Corpus Data","desc":"Data content"},{"title":"Data size","content":"7.29 million pairs of Chinese-Vietnamese Parallel Corpus Data","desc":"Data size"},{"title":"Language","content":"Chinese,Vietnamese","desc":"Language"},{"title":"Application scenario","content":"machine translation","desc":"Application scenario"},{"title":"Accuracy rate","content":"90%","desc":"Accuracy rate"}],"datatag":"Chinese,Vietnamese,Chinese-Vietnamese,Parallel Corpus","technologydoc":null,"downurl":null,"datainfo":null,"standard":null,"dataylurl":null,"flag":null,"publishtime":null,"createby":null,"createtime":null,"ext1":null,"samplestoreloc":null,"hosturl":null,"datasize":null,"industryPlan":null,"keyInformation":"","samplePresentation":[{"name":"/data/apps/damp/temp/ziptemp/APY220317001_demo1711015208648/zh-vi ????.png","url":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY220317001_demo1711015208648/zh-vi%20%3F%3F%3F%3F.png?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=6bJid711OD71%2BZiotlLECOimk1I%3D","intro":"","size":0,"progress":100,"type":"jpg"}],"officialSummary":"This dataset contains 7.29 million Chinese-Vietnamese parallel sentence pairs stored in TXT format. The corpus covers multiple domains, including tourism, healthcare, daily life, news, and other topics. The data has undergone cleaning, anonymization, and quality inspection to improve data quality and protect sensitive information. It can be used as a basic corpus for text data analysis in fields such as machine translation.","dataexampl":null,"datakeyword":["Chinese Vietnamese parallel corpus","Chinese Vietnamese translation dataset","Chinese Vietnamese parallel dataset","Chinese Vietnamese sentence pairs","Chinese Vietnamese translation corpus"],"isDelete":null,"ids":null,"idsList":null,"datasetCode":null,"productStatus":null,"tagTypeEn":"Type","tagTypeZh":null,"website":null,"samplePresentationList":null,"datazyList":null,"keyInformationList":null,"dataexamplList":null,"bgimg":null,"datazyScriptList":null,"datakeywordListString":null,"sourceShowPage":"nlu","dataShowType":"[{\"code\":\"0\",\"language\":\"ZH\"},{\"code\":\"1\",\"language\":\"ZH\"},{\"code\":\"2\",\"language\":\"EN,PT,DE,KO,FR,ES\"},{\"code\":\"3\",\"language\":\"EN\"},{\"code\":\"4\",\"language\":\"JP\"}]","productNameEn":"7,290,000 Groups -Chinese -Vietnamese Parallel Corpus Data","BGimg":"","voiceBg":["/shujutang/static/image/comm/audio_bg.webp","/shujutang/static/image/comm/audio_bg2.webp","/shujutang/static/image/comm/audio_bg3.webp","/shujutang/static/image/comm/audio_bg4.webp","/shujutang/static/image/comm/audio_bg5.webp"]}
https://www.nexdata.ai/shujutang/static/image/index/datatang_wenben_default.webp
[{"@type":"ImageObject","embedUrl":"https://bj-oss-datatang-03.oss-cn-beijing.aliyuncs.com/filesInfoUpload/data/apps/damp/temp/ziptemp/APY220317001_demo1711015208648/zh-vi%20%3F%3F%3F%3F.png?Expires=4102329599&OSSAccessKeyId=LTAI8NWs2pDolLNH&Signature=6bJid711OD71%2BZiotlLECOimk1I%3D"}]
7.29M Chinese-Vietnamese Sentence Pairs – Machine Translation Dataset
Chinese Vietnamese parallel corpus
Chinese Vietnamese translation dataset
Chinese Vietnamese parallel dataset
Chinese Vietnamese sentence pairs
Chinese Vietnamese translation corpus
This dataset contains 7.29 million Chinese-Vietnamese parallel sentence pairs stored in TXT format. The corpus covers multiple domains, including tourism, healthcare, daily life, news, and other topics. The data has undergone cleaning, anonymization, and quality inspection to improve data quality and protect sensitive information. It can be used as a basic corpus for text data analysis in fields such as machine translation.
This is a paid datasets for commercial use, research purpose and more. Licensed ready made datasets help jump-start AI projects.
![Specifications]()
Specifications
Data content
Chinese-Vietnamese Parallel Corpus Data
Data size
7.29 million pairs of Chinese-Vietnamese Parallel Corpus Data
Language
Chinese,Vietnamese
Application scenario
machine translation
![Sample]()
Sample
![Recommended Datasets]()
Recommended Dataset
Tell Us Your Special Needs
a9cd0421-fe6b-411c-8907-8733f35148c2