Modul:hi-translit
Penampilan
- Berikut merupakan pendokumenan yang dijana oleh Modul:pendokumenan/functions/translit. [sunting]
- Pautan berguna: senarai sublaman • pautan • transklusi • kes ujian • kotak pasir
Modul ini akan mentransliterasi Bahasa Hindi teks. Ia juga digunakan untuk mentransliterasi bahasa Angika, Awadhi, Bagheli, Mahasu Pahari, Haryanvi, Bhili, Bhadrawahi, Bundeli, Palya Bareli, Braj, Chambeali, Churahi, Dogri, Gaddi, Garhwali, Chhattisgarhi, Bilaspuri, Kullu Pahari, Kurukh, Mandeali, Mewari, Malvi, Marwari, Nimadi, Pangwali, Mundari, and Kangri.
The module should preferably not be called directly from templates or other modules.
To use it from a template, use {{xlit}}.
Within a module, use Module:languages#Language:transliterate.
For testcases, see Module:hi-translit/testcases.
Functions
[sunting]tr(text, lang, sc)- Transliterates a given piece of
textwritten in the script specified by the codesc, and language specified by the codelang. - When the transliteration fails, returns
nil.
-- Transliteration for Hindi (possibly other languages using Devanagari script, except for Sanskrit)
local export = {}
local gsub = mw.ustring.gsub
local match = mw.ustring.match
local conv = {
-- consonants
['क'] = 'k', ['ख'] = 'kh', ['ग'] = 'g', ['घ'] = 'gh', ['ङ'] = 'ṅ',
['च'] = 'c', ['छ'] = 'ch', ['ज'] = 'j', ['झ'] = 'jh', ['ञ'] = 'ñ',
['ट'] = 'ṭ', ['ठ'] = 'ṭh', ['ड'] = 'ḍ', ['ढ'] = 'ḍh', ['ण'] = 'ṇ',
['त'] = 't', ['थ'] = 'th', ['द'] = 'd', ['ध'] = 'dh', ['न'] = 'n',
['प'] = 'p', ['फ'] = 'ph', ['ब'] = 'b', ['भ'] = 'bh', ['म'] = 'm',
['य'] = 'y', ['र'] = 'r', ['ल'] = 'l', ['व'] = 'v', ['ळ'] = 'ḷ',
['श'] = 'ś', ['ष'] = 'ṣ', ['स'] = 's', ['ह'] = 'h',
['क़'] = 'q', ['ख़'] = 'x', ['ग़'] = 'ġ', ['ऴ'] = 'ḻ',
['ज़'] = 'z', ['झ़'] = 'ž', ['ड़'] = 'ṛ', ['ढ़'] = 'ṛh',
['फ़'] = 'f', ['थ़'] = 'θ', ['ऩ'] = 'ṉ', ['ऱ'] = 'ṟ',
-- vowel diacritics
['ि'] = 'i', ['ु'] = 'u', ['े'] = 'e', ['ो'] = 'o',
['ा'] = 'ā', ['ी'] = 'ī', ['ू'] = 'ū',
['ृ'] = 'ŕ',
['ै'] = 'ai', ['ौ'] = 'au',
-- vowel signs
['अ'] = 'a', ['इ'] = 'i', ['उ'] = 'u', ['ए'] = 'e', ['ओ'] = 'o',
['आ'] = 'ā', ['ई'] = 'ī', ['ऊ'] = 'ū',
['ऋ'] = 'ŕ',
['ऐ'] = 'ai', ['औ'] = 'au',
-- chandrabindu
['ँ'] = '̃',
-- anusvara
['ं'] = 'ṁ',
-- visarga
['ः'] = 'ḥ',
-- virama
['्'] = '',
-- numerals
['०'] = '0', ['१'] = '1', ['२'] = '2', ['३'] = '3', ['४'] = '4', ['५'] = '5', ['६'] = '6', ['७'] = '7', ['८'] = '8', ['९'] = '9',
-- punctuation
['।'] = '.', -- danda
['+'] = '', -- compound separator
}
local nasal_assim = {
['क'] = 'ङ', ['ख'] = 'ङ', ['ग'] = 'ङ', ['घ'] = 'ङ',
['च'] = 'ञ', ['छ'] = 'ञ', ['ज'] = 'ञ', ['झ'] = 'ञ',
['ट'] = 'ण', ['ठ'] = 'ण', ['ड'] = 'ण', ['ढ'] = 'ण',
['प'] = 'म', ['फ'] = 'म', ['ब'] = 'म', ['भ'] = 'म', ['म'] = 'म',
}
local all_cons, special_cons = 'कखगघङचछजझञटठडढतथदधपफबभशषसयरलवहणनम', 'यरलवहनम'
local vowel, vowel_sign = 'aिुृेोाीूैौ', 'अइउएओआईऊऋऐऔ'
local syncope_pattern = '([' .. vowel .. vowel_sign .. '])([' .. all_cons .. '])a([' .. gsub(all_cons, "य", "") .. '])(ं?[' .. vowel .. vowel_sign .. '])'
local function rev_string(text)
local result, length = '', mw.ustring.len(text)
for i = 1, length do
result = result .. mw.ustring.sub(text, length - i + 1, length - i + 1)
end
return result
end
function export.tr(text, lang, sc)
text = gsub(text, '([' .. all_cons .. ']़?)([' .. vowel .. '्]?)', function(c, d)
return c .. (d == "" and 'a' or d) end)
local result = {}
for word in mw.text.gsplit(text, " ", true) do
word = rev_string(word)
word = gsub(word, '^a(़?)([' .. all_cons .. '])(.)', function(opt, first, second)
return (((match(first, '[' .. special_cons .. ']') and match(second, '्')) or match(first .. second, 'य[ीेै]'))
and 'a' or "") .. opt .. first .. second end)
while match(word, syncope_pattern) do
word = gsub(word, syncope_pattern, '%1%2%3%4')
end
word = gsub(word, '(.?)ं(.)', function(succ, prev)
return succ .. (succ..prev == "a" and "्म" or
(succ == "" and match(prev, '[' .. vowel .. ']') and "̃" or nasal_assim[succ] or "n")) .. prev end)
table.insert(result, rev_string(word))
end
text = table.concat(result, " ")
text = gsub(text, '.़?', conv)
return mw.ustring.toNFC(text)
end
return export
Kategori:
- Modul transliterasi digunakan oleh 28 bahasa
- Modul bahasa Hindi
- Modul transliterasi
- Modul bahasa Haryanvi
- Modul bahasa Gaddi
- Modul bahasa Angika
- Modul bahasa Dogri
- Modul bahasa Kurukh
- Modul bahasa Marwari
- Modul bahasa Mundari
- Modul bahasa Mandeali
- Modul bahasa Bundeli
- Modul bahasa Braj
- Modul bahasa Bagheli
- Modul bahasa Awadhi
- Modul bahasa Palya Bareli
- Modul bahasa Nimadi
- Modul bahasa Chambeali
- Modul bahasa Kullu Pahari
- Modul bahasa Kangri
- Modul bahasa Garhwali
- Modul bahasa Bilaspuri
- Modul bahasa Bhadrawahi
- Modul bahasa Bhili
- Modul bahasa Chhattisgarhi
- Modul bahasa Pangwali
- Modul bahasa Malvi
- Modul bahasa Churahi
- Modul bahasa Mahasu Pahari
- Modul bahasa Mewari