feat(i18n): offer every ISO 639-1 language

The table was thirty languages typed by hand. It is now generated:
utils/generate_languages.py takes the 184 alpha-2 codes from pycountry,
the display names from CLDR through babel, and the writing direction from
CLDR character order. 159 languages carry their endonym; the remaining 25
have no CLDR entry and carry their English ISO name.

That corrects the right-to-left set, which had four entries and needs ten
— dv, ks, ps, sd, ug and yi were simply missed.

Only 29 languages ship an interface catalogue, so the other 155 render in
English until one is filled. make i18n-ui fills app/i18n/ui/ for them, and
make i18n now covers the interface strings as well; neither asks for a
string a shipped catalogue already answers, so hand-written entries stay.

184 entries do not fit on a screen, so the language menu scrolls inside
itself. overscroll-behavior keeps the page behind it from moving once the
list reaches its end.

babel and pycountry are dev dependencies: the generator needs them, the
application does not.

Co-Authored-By: Claude Opus 5 (1M context) <noreply@anthropic.com>
This commit is contained in:
2026-08-22 01:08:28 +02:00
parent 0c67f999f6
commit ef1c8ff09a
12 changed files with 413 additions and 76 deletions

196
app/utils/languages.py Normal file
View File

@@ -0,0 +1,196 @@
"""ISO 639-1 languages, their display names and writing direction.
Generated by utils/generate_languages.py — edit that script, not this file.
Display names are CLDR endonyms where CLDR covers the language and the English
ISO 639-1 name for the 25 codes it does not.
"""
LANGUAGES = {
"en": "English",
"aa": "Qafar",
"ab": "Аԥсшәа",
"ae": "Avestan",
"af": "Afrikaans",
"ak": "Akan",
"am": "አማርኛ",
"an": "aragonés",
"ar": "العربية",
"as": "অসমীয়া",
"av": "Avaric",
"ay": "Aymara",
"az": "azərbaycan",
"ba": "башҡорт теле",
"be": "беларуская",
"bg": "български",
"bi": "Bislama",
"bm": "bamanakan",
"bn": "বাংলা",
"bo": "བོད་སྐད་",
"br": "brezhoneg",
"bs": "bosanski",
"ca": "català",
"ce": "нохчийн",
"ch": "Chamorro",
"co": "corsu",
"cr": "Cree",
"cs": "čeština",
"cu": "церковнослове́нскїй",
"cv": "чӑваш",
"cy": "Cymraeg",
"da": "dansk",
"de": "Deutsch",
"dv": "ދިވެހިބަސް",
"dz": "རྫོང་ཁ",
"ee": "eʋegbe",
"el": "Ελληνικά",
"eo": "Esperanto",
"es": "español",
"et": "eesti",
"eu": "euskara",
"fa": "فارسی",
"ff": "Pulaar",
"fi": "suomi",
"fj": "Fijian",
"fo": "føroyskt",
"fr": "français",
"fy": "Frysk",
"ga": "Gaeilge",
"gd": "Gàidhlig",
"gl": "galego",
"gn": "avañe",
"gu": "ગુજરાતી",
"gv": "Gaelg",
"ha": "Hausa",
"he": "עברית",
"hi": "हिन्दी",
"ho": "Hiri Motu",
"hr": "hrvatski",
"ht": "Kreyòl Ayisyen",
"hu": "magyar",
"hy": "հայերեն",
"hz": "Herero",
"ia": "interlingua",
"id": "Indonesia",
"ie": "Interlingue",
"ig": "Igbo",
"ii": "ꆈꌠꉙ",
"ik": "Inupiaq",
"io": "Ido",
"is": "íslenska",
"it": "italiano",
"iu": "ᐃᓄᒃᑎᑐᑦ",
"ja": "日本語",
"jv": "Jawa",
"ka": "ქართული",
"kg": "Kongo",
"ki": "Gikuyu",
"kj": "Kuanyama",
"kk": "қазақ тілі",
"kl": "kalaallisut",
"km": "ខ្មែរ",
"kn": "ಕನ್ನಡ",
"ko": "한국어",
"kr": "Kanuri",
"ks": "کٲشُر",
"ku": "kurdî (kurmancî)",
"kv": "Komi",
"kw": "kernewek",
"ky": "кыргызча",
"la": "Latina",
"lb": "Lëtzebuergesch",
"lg": "Luganda",
"li": "Limburgan",
"ln": "lingála",
"lo": "ລາວ",
"lt": "lietuvių",
"lu": "Tshiluba",
"lv": "latviešu",
"mg": "Malagasy",
"mh": "Marshallese",
"mi": "Māori",
"mk": "македонски",
"ml": "മലയാളം",
"mn": "монгол",
"mr": "मराठी",
"ms": "Melayu",
"mt": "Malti",
"my": "မြန်မာ",
"na": "Nauru",
"nb": "norsk bokmål",
"nd": "isiNdebele",
"ne": "नेपाली",
"ng": "Ndonga",
"nl": "Nederlands",
"nn": "norsk nynorsk",
"no": "norsk",
"nr": "isiNdebele",
"nv": "Diné Bizaad",
"ny": "Nyanja",
"oc": "occitan",
"oj": "Ojibwa",
"om": "Oromoo",
"or": "ଓଡ଼ିଆ",
"os": "ирон",
"pa": "ਪੰਜਾਬੀ",
"pi": "Pali",
"pl": "polski",
"ps": "پښتو",
"pt": "português",
"qu": "Runasimi",
"rm": "rumantsch",
"rn": "Ikirundi",
"ro": "română",
"ru": "русский",
"rw": "Ikinyarwanda",
"sa": "संस्कृत भाषा",
"sc": "sardu",
"sd": "سنڌي",
"se": "davvisámegiella",
"sg": "Sängö",
"sh": "Serbo-Croatian",
"si": "සිංහල",
"sk": "slovenčina",
"sl": "slovenščina",
"sm": "Samoan",
"sn": "chiShona",
"so": "Soomaali",
"sq": "shqip",
"sr": "српски",
"ss": "siSwati",
"st": "Sesotho",
"su": "Basa Sunda",
"sv": "svenska",
"sw": "Kiswahili",
"ta": "தமிழ்",
"te": "తెలుగు",
"tg": "тоҷикӣ",
"th": "ไทย",
"ti": "ትግርኛ",
"tk": "türkmen dili",
"tl": "Tagalog",
"tn": "Setswana",
"to": "lea fakatonga",
"tr": "Türkçe",
"ts": "Xitsonga",
"tt": "татар",
"tw": "Twi",
"ty": "Tahitian",
"ug": "ئۇيغۇرچە",
"uk": "українська",
"ur": "اردو",
"uz": "ozbek",
"ve": "Tshivenḓa",
"vi": "Tiếng Việt",
"vo": "Volapük",
"wa": "walon",
"wo": "Wolof",
"xh": "IsiXhosa",
"yi": "ייִדיש",
"yo": "Èdè Yorùbá",
"za": "Vahcuengh",
"zh": "中文",
"zu": "isiZulu",
}
RTL_LANGUAGES = frozenset({"ar", "dv", "fa", "he", "ks", "ps", "sd", "ug", "ur", "yi"})