summaryrefslogtreecommitdiffstats
path: root/util/local_database/cldr2qlocalexml.py
diff options
context:
space:
mode:
Diffstat (limited to 'util/local_database/cldr2qlocalexml.py')
-rwxr-xr-xutil/local_database/cldr2qlocalexml.py459
1 files changed, 459 insertions, 0 deletions
diff --git a/util/local_database/cldr2qlocalexml.py b/util/local_database/cldr2qlocalexml.py
new file mode 100755
index 0000000..e4446c4
--- /dev/null
+++ b/util/local_database/cldr2qlocalexml.py
@@ -0,0 +1,459 @@
+#! /usr/bin/python
+
+import os
+import sys
+import enumdata
+import xpathlite
+from xpathlite import DraftResolution
+import re
+
+findEntry = xpathlite.findEntry
+
+def ordStr(c):
+ if len(c) == 1:
+ return str(ord(c))
+ return "##########"
+
+# the following functions are supposed to fix the problem with QLocale
+# returning a character instead of strings for QLocale::exponential()
+# and others. So we fallback to default values in these cases.
+def fixOrdStrExp(c):
+ if len(c) == 1:
+ return str(ord(c))
+ return str(ord('e'))
+def fixOrdStrPercent(c):
+ if len(c) == 1:
+ return str(ord(c))
+ return str(ord('%'))
+def fixOrdStrList(c):
+ if len(c) == 1:
+ return str(ord(c))
+ return str(ord(';'))
+
+def generateLocaleInfo(path):
+ (dir_name, file_name) = os.path.split(path)
+
+ exp = re.compile(r"([a-z]+)_([A-Z]{2})\.xml")
+ m = exp.match(file_name)
+ if not m:
+ return {}
+
+ language_code = m.group(1)
+ country_code = m.group(2)
+
+ language_id = enumdata.languageCodeToId(language_code)
+ if language_id == -1:
+ sys.stderr.write("unnknown language code \"" + language_code + "\"\n")
+ return {}
+ language = enumdata.language_list[language_id][0]
+
+ country_id = enumdata.countryCodeToId(country_code)
+ if country_id == -1:
+ sys.stderr.write("unnknown country code \"" + country_code + "\"\n")
+ return {}
+ country = enumdata.country_list[country_id][0]
+
+ base = dir_name + "/" + language_code + "_" + country_code
+
+ result = {}
+ result['base'] = base
+
+ result['language'] = language
+ result['country'] = country
+ result['language_id'] = language_id
+ result['country_id'] = country_id
+ result['decimal'] = findEntry(base, "numbers/symbols/decimal")
+ result['group'] = findEntry(base, "numbers/symbols/group")
+ result['list'] = findEntry(base, "numbers/symbols/list")
+ result['percent'] = findEntry(base, "numbers/symbols/percentSign")
+ result['zero'] = findEntry(base, "numbers/symbols/nativeZeroDigit")
+ result['minus'] = findEntry(base, "numbers/symbols/minusSign")
+ result['plus'] = findEntry(base, "numbers/symbols/plusSign")
+ result['exp'] = findEntry(base, "numbers/symbols/exponential").lower()
+ result['am'] = findEntry(base, "dates/calendars/calendar[gregorian]/am", draft=DraftResolution.approved)
+ result['pm'] = findEntry(base, "dates/calendars/calendar[gregorian]/pm", draft=DraftResolution.approved)
+ result['longDateFormat'] = findEntry(base, "dates/calendars/calendar[gregorian]/dateFormats/dateFormatLength[full]/dateFormat/pattern")
+ result['shortDateFormat'] = findEntry(base, "dates/calendars/calendar[gregorian]/dateFormats/dateFormatLength[short]/dateFormat/pattern")
+ result['longTimeFormat'] = findEntry(base, "dates/calendars/calendar[gregorian]/timeFormats/timeFormatLength[full]/timeFormat/pattern")
+ result['shortTimeFormat'] = findEntry(base, "dates/calendars/calendar[gregorian]/timeFormats/timeFormatLength[short]/timeFormat/pattern")
+
+ standalone_long_month_path = "dates/calendars/calendar[gregorian]/months/monthContext[stand-alone]/monthWidth[wide]/month"
+ result['standaloneLongMonths'] \
+ = findEntry(base, standalone_long_month_path + "[1]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[2]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[3]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[4]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[5]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[6]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[7]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[8]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[9]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[10]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[11]") + ";" \
+ + findEntry(base, standalone_long_month_path + "[12]") + ";"
+
+ standalone_short_month_path = "dates/calendars/calendar[gregorian]/months/monthContext[stand-alone]/monthWidth[abbreviated]/month"
+ result['standaloneShortMonths'] \
+ = findEntry(base, standalone_short_month_path + "[1]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[2]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[3]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[4]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[5]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[6]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[7]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[8]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[9]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[10]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[11]") + ";" \
+ + findEntry(base, standalone_short_month_path + "[12]") + ";"
+
+ standalone_narrow_month_path = "dates/calendars/calendar[gregorian]/months/monthContext[stand-alone]/monthWidth[narrow]/month"
+ result['standaloneNarrowMonths'] \
+ = findEntry(base, standalone_narrow_month_path + "[1]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[2]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[3]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[4]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[5]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[6]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[7]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[8]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[9]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[10]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[11]") + ";" \
+ + findEntry(base, standalone_narrow_month_path + "[12]") + ";"
+
+ long_month_path = "dates/calendars/calendar[gregorian]/months/monthContext[format]/monthWidth[wide]/month"
+ result['longMonths'] \
+ = findEntry(base, long_month_path + "[1]") + ";" \
+ + findEntry(base, long_month_path + "[2]") + ";" \
+ + findEntry(base, long_month_path + "[3]") + ";" \
+ + findEntry(base, long_month_path + "[4]") + ";" \
+ + findEntry(base, long_month_path + "[5]") + ";" \
+ + findEntry(base, long_month_path + "[6]") + ";" \
+ + findEntry(base, long_month_path + "[7]") + ";" \
+ + findEntry(base, long_month_path + "[8]") + ";" \
+ + findEntry(base, long_month_path + "[9]") + ";" \
+ + findEntry(base, long_month_path + "[10]") + ";" \
+ + findEntry(base, long_month_path + "[11]") + ";" \
+ + findEntry(base, long_month_path + "[12]") + ";"
+
+ short_month_path = "dates/calendars/calendar[gregorian]/months/monthContext[format]/monthWidth[abbreviated]/month"
+ result['shortMonths'] \
+ = findEntry(base, short_month_path + "[1]") + ";" \
+ + findEntry(base, short_month_path + "[2]") + ";" \
+ + findEntry(base, short_month_path + "[3]") + ";" \
+ + findEntry(base, short_month_path + "[4]") + ";" \
+ + findEntry(base, short_month_path + "[5]") + ";" \
+ + findEntry(base, short_month_path + "[6]") + ";" \
+ + findEntry(base, short_month_path + "[7]") + ";" \
+ + findEntry(base, short_month_path + "[8]") + ";" \
+ + findEntry(base, short_month_path + "[9]") + ";" \
+ + findEntry(base, short_month_path + "[10]") + ";" \
+ + findEntry(base, short_month_path + "[11]") + ";" \
+ + findEntry(base, short_month_path + "[12]") + ";"
+
+ narrow_month_path = "dates/calendars/calendar[gregorian]/months/monthContext[format]/monthWidth[narrow]/month"
+ result['narrowMonths'] \
+ = findEntry(base, narrow_month_path + "[1]") + ";" \
+ + findEntry(base, narrow_month_path + "[2]") + ";" \
+ + findEntry(base, narrow_month_path + "[3]") + ";" \
+ + findEntry(base, narrow_month_path + "[4]") + ";" \
+ + findEntry(base, narrow_month_path + "[5]") + ";" \
+ + findEntry(base, narrow_month_path + "[6]") + ";" \
+ + findEntry(base, narrow_month_path + "[7]") + ";" \
+ + findEntry(base, narrow_month_path + "[8]") + ";" \
+ + findEntry(base, narrow_month_path + "[9]") + ";" \
+ + findEntry(base, narrow_month_path + "[10]") + ";" \
+ + findEntry(base, narrow_month_path + "[11]") + ";" \
+ + findEntry(base, narrow_month_path + "[12]") + ";"
+
+ long_day_path = "dates/calendars/calendar[gregorian]/days/dayContext[format]/dayWidth[wide]/day"
+ result['longDays'] \
+ = findEntry(base, long_day_path + "[sun]") + ";" \
+ + findEntry(base, long_day_path + "[mon]") + ";" \
+ + findEntry(base, long_day_path + "[tue]") + ";" \
+ + findEntry(base, long_day_path + "[wed]") + ";" \
+ + findEntry(base, long_day_path + "[thu]") + ";" \
+ + findEntry(base, long_day_path + "[fri]") + ";" \
+ + findEntry(base, long_day_path + "[sat]") + ";"
+
+ short_day_path = "dates/calendars/calendar[gregorian]/days/dayContext[format]/dayWidth[abbreviated]/day"
+ result['shortDays'] \
+ = findEntry(base, short_day_path + "[sun]") + ";" \
+ + findEntry(base, short_day_path + "[mon]") + ";" \
+ + findEntry(base, short_day_path + "[tue]") + ";" \
+ + findEntry(base, short_day_path + "[wed]") + ";" \
+ + findEntry(base, short_day_path + "[thu]") + ";" \
+ + findEntry(base, short_day_path + "[fri]") + ";" \
+ + findEntry(base, short_day_path + "[sat]") + ";"
+
+ narrow_day_path = "dates/calendars/calendar[gregorian]/days/dayContext[format]/dayWidth[narrow]/day"
+ result['narrowDays'] \
+ = findEntry(base, narrow_day_path + "[sun]") + ";" \
+ + findEntry(base, narrow_day_path + "[mon]") + ";" \
+ + findEntry(base, narrow_day_path + "[tue]") + ";" \
+ + findEntry(base, narrow_day_path + "[wed]") + ";" \
+ + findEntry(base, narrow_day_path + "[thu]") + ";" \
+ + findEntry(base, narrow_day_path + "[fri]") + ";" \
+ + findEntry(base, narrow_day_path + "[sat]") + ";"
+
+ standalone_long_day_path = "dates/calendars/calendar[gregorian]/days/dayContext[stand-alone]/dayWidth[wide]/day"
+ result['standaloneLongDays'] \
+ = findEntry(base, standalone_long_day_path + "[sun]") + ";" \
+ + findEntry(base, standalone_long_day_path + "[mon]") + ";" \
+ + findEntry(base, standalone_long_day_path + "[tue]") + ";" \
+ + findEntry(base, standalone_long_day_path + "[wed]") + ";" \
+ + findEntry(base, standalone_long_day_path + "[thu]") + ";" \
+ + findEntry(base, standalone_long_day_path + "[fri]") + ";" \
+ + findEntry(base, standalone_long_day_path + "[sat]") + ";"
+
+ standalone_short_day_path = "dates/calendars/calendar[gregorian]/days/dayContext[stand-alone]/dayWidth[abbreviated]/day"
+ result['standaloneShortDays'] \
+ = findEntry(base, standalone_short_day_path + "[sun]") + ";" \
+ + findEntry(base, standalone_short_day_path + "[mon]") + ";" \
+ + findEntry(base, standalone_short_day_path + "[tue]") + ";" \
+ + findEntry(base, standalone_short_day_path + "[wed]") + ";" \
+ + findEntry(base, standalone_short_day_path + "[thu]") + ";" \
+ + findEntry(base, standalone_short_day_path + "[fri]") + ";" \
+ + findEntry(base, standalone_short_day_path + "[sat]") + ";"
+
+ standalone_narrow_day_path = "dates/calendars/calendar[gregorian]/days/dayContext[stand-alone]/dayWidth[narrow]/day"
+ result['standaloneNarrowDays'] \
+ = findEntry(base, standalone_narrow_day_path + "[sun]") + ";" \
+ + findEntry(base, standalone_narrow_day_path + "[mon]") + ";" \
+ + findEntry(base, standalone_narrow_day_path + "[tue]") + ";" \
+ + findEntry(base, standalone_narrow_day_path + "[wed]") + ";" \
+ + findEntry(base, standalone_narrow_day_path + "[thu]") + ";" \
+ + findEntry(base, standalone_narrow_day_path + "[fri]") + ";" \
+ + findEntry(base, standalone_narrow_day_path + "[sat]") + ";"
+
+
+ return result
+
+def addEscapes(s):
+ result = ''
+ for c in s:
+ n = ord(c)
+ if n < 128:
+ result += c
+ else:
+ result += "\\x"
+ result += "%02x" % (n)
+ return result
+
+def unicodeStr(s):
+ utf8 = s.encode('utf-8')
+ return "<size>" + str(len(utf8)) + "</size><data>" + addEscapes(utf8) + "</data>"
+
+def usage():
+ print "Usage: cldr2qlocalexml.py <path-to-cldr-main>"
+ sys.exit()
+
+if len(sys.argv) != 2:
+ usage()
+
+cldr_dir = sys.argv[1]
+
+if not os.path.isdir(cldr_dir):
+ usage()
+
+cldr_files = os.listdir(cldr_dir)
+
+locale_database = {}
+for file in cldr_files:
+ l = generateLocaleInfo(cldr_dir + "/" + file)
+ if not l:
+ sys.stderr.write("skipping file \"" + file + "\"\n")
+ continue
+
+ locale_database[(l['language_id'], l['country_id'])] = l
+
+locale_keys = locale_database.keys()
+locale_keys.sort()
+
+print "<localeDatabase>"
+print " <languageList>"
+for id in enumdata.language_list:
+ l = enumdata.language_list[id]
+ print " <language>"
+ print " <name>" + l[0] + "</name>"
+ print " <id>" + str(id) + "</id>"
+ print " <code>" + l[1] + "</code>"
+ print " </language>"
+print " </languageList>"
+
+print " <countryList>"
+for id in enumdata.country_list:
+ l = enumdata.country_list[id]
+ print " <country>"
+ print " <name>" + l[0] + "</name>"
+ print " <id>" + str(id) + "</id>"
+ print " <code>" + l[1] + "</code>"
+ print " </country>"
+print " </countryList>"
+
+print \
+" <defaultCountryList>\n\
+ <defaultCountry>\n\
+ <language>Afrikaans</language>\n\
+ <country>SouthAfrica</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Afan</language>\n\
+ <country>Ethiopia</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Afar</language>\n\
+ <country>Djibouti</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Arabic</language>\n\
+ <country>SaudiArabia</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Chinese</language>\n\
+ <country>China</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Dutch</language>\n\
+ <country>Netherlands</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>English</language>\n\
+ <country>UnitedStates</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>French</language>\n\
+ <country>France</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>German</language>\n\
+ <country>Germany</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Greek</language>\n\
+ <country>Greece</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Italian</language>\n\
+ <country>Italy</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Malay</language>\n\
+ <country>Malaysia</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Portuguese</language>\n\
+ <country>Portugal</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Russian</language>\n\
+ <country>RussianFederation</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Serbian</language>\n\
+ <country>SerbiaAndMontenegro</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>SerboCroatian</language>\n\
+ <country>SerbiaAndMontenegro</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Somali</language>\n\
+ <country>Somalia</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Spanish</language>\n\
+ <country>Spain</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Swahili</language>\n\
+ <country>Kenya</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Swedish</language>\n\
+ <country>Sweden</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Tigrinya</language>\n\
+ <country>Eritrea</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Uzbek</language>\n\
+ <country>Uzbekistan</country>\n\
+ </defaultCountry>\n\
+ <defaultCountry>\n\
+ <language>Persian</language>\n\
+ <country>Iran</country>\n\
+ </defaultCountry>\n\
+ </defaultCountryList>"
+
+print " <localeList>"
+print \
+" <locale>\n\
+ <language>C</language>\n\
+ <country>AnyCountry</country>\n\
+ <decimal>46</decimal>\n\
+ <group>44</group>\n\
+ <list>59</list>\n\
+ <percent>37</percent>\n\
+ <zero>48</zero>\n\
+ <minus>45</minus>\n\
+ <plus>43</plus>\n\
+ <exp>101</exp>\n\
+ <am>AM</am>\n\
+ <pm>PM</pm>\n\
+ <longDateFormat>EEEE, d MMMM yyyy</longDateFormat>\n\
+ <shortDateFormat>d MMM yyyy</shortDateFormat>\n\
+ <longTimeFormat>HH:mm:ss z</longTimeFormat>\n\
+ <shortTimeFormat>HH:mm:ss</shortTimeFormat>\n\
+ <standaloneLongMonths>January;February;March;April;May;June;July;August;September;October;November;December;</standaloneLongMonths>\n\
+ <standaloneShortMonths>Jan;Feb;Mar;Apr;May;Jun;Jul;Aug;Sep;Oct;Nov;Dec;</standaloneShortMonths>\n\
+ <standaloneNarrowMonths>J;F;M;A;M;J;J;A;S;O;N;D;</standaloneNarrowMonths>\n\
+ <longMonths>January;February;March;April;May;June;July;August;September;October;November;December;</longMonths>\n\
+ <shortMonths>Jan;Feb;Mar;Apr;May;Jun;Jul;Aug;Sep;Oct;Nov;Dec;</shortMonths>\n\
+ <narrowMonths>1;2;3;4;5;6;7;8;9;10;11;12;</narrowMonths>\n\
+ <longDays>Sunday;Monday;Tuesday;Wednesday;Thursday;Friday;Saturday;</longDays>\n\
+ <shortDays>Sun;Mon;Tue;Wed;Thu;Fri;Sat;</shortDays>\n\
+ <narrowDays>7;1;2;3;4;5;6;</narrowDays>\n\
+ <standaloneLongDays>Sunday;Monday;Tuesday;Wednesday;Thursday;Friday;Saturday;</standaloneLongDays>\n\
+ <standaloneShortDays>Sun;Mon;Tue;Wed;Thu;Fri;Sat;</standaloneShortDays>\n\
+ <standaloneNarrowDays>S;M;T;W;T;F;S;</standaloneNarrowDays>\n\
+ </locale>"
+
+for key in locale_keys:
+ l = locale_database[key]
+
+ print " <locale>"
+# print " <source>" + l['base'] + "</source>"
+ print " <language>" + l['language'] + "</language>"
+ print " <country>" + l['country'] + "</country>"
+ print " <decimal>" + ordStr(l['decimal']) + "</decimal>"
+ print " <group>" + ordStr(l['group']) + "</group>"
+ print " <list>" + fixOrdStrList(l['list']) + "</list>"
+ print " <percent>" + fixOrdStrPercent(l['percent']) + "</percent>"
+ print " <zero>" + ordStr(l['zero']) + "</zero>"
+ print " <minus>" + ordStr(l['minus']) + "</minus>"
+ print " <plus>" + ordStr(l['plus']) + "</plus>"
+ print " <exp>" + fixOrdStrExp(l['exp']) + "</exp>"
+ print " <am>" + l['am'].encode('utf-8') + "</am>"
+ print " <pm>" + l['pm'].encode('utf-8') + "</pm>"
+ print " <longDateFormat>" + l['longDateFormat'].encode('utf-8') + "</longDateFormat>"
+ print " <shortDateFormat>" + l['shortDateFormat'].encode('utf-8') + "</shortDateFormat>"
+ print " <longTimeFormat>" + l['longTimeFormat'].encode('utf-8') + "</longTimeFormat>"
+ print " <shortTimeFormat>" + l['shortTimeFormat'].encode('utf-8') + "</shortTimeFormat>"
+ print " <standaloneLongMonths>" + l['standaloneLongMonths'].encode('utf-8') + "</standaloneLongMonths>"
+ print " <standaloneShortMonths>"+ l['standaloneShortMonths'].encode('utf-8') + "</standaloneShortMonths>"
+ print " <standaloneNarrowMonths>"+ l['standaloneNarrowMonths'].encode('utf-8') + "</standaloneNarrowMonths>"
+ print " <longMonths>" + l['longMonths'].encode('utf-8') + "</longMonths>"
+ print " <shortMonths>" + l['shortMonths'].encode('utf-8') + "</shortMonths>"
+ print " <narrowMonths>" + l['narrowMonths'].encode('utf-8') + "</narrowMonths>"
+ print " <longDays>" + l['longDays'].encode('utf-8') + "</longDays>"
+ print " <shortDays>" + l['shortDays'].encode('utf-8') + "</shortDays>"
+ print " <narrowDays>" + l['narrowDays'].encode('utf-8') + "</narrowDays>"
+ print " <standaloneLongDays>" + l['standaloneLongDays'].encode('utf-8') + "</standaloneLongDays>"
+ print " <standaloneShortDays>" + l['standaloneShortDays'].encode('utf-8') + "</standaloneShortDays>"
+ print " <standaloneNarrowDays>" + l['standaloneNarrowDays'].encode('utf-8') + "</standaloneNarrowDays>"
+ print " </locale>"
+print " </localeList>"
+print "</localeDatabase>"