searx/searx/engines/lingva.py

# SPDX-License-Identifier: AGPL-3.0-or-later
# lint: pylint
"""Lingva (alternative Google Translate frontend)"""

from json import loads

about = {
    "website": 'https://lingva.ml',
    "wikidata_id": None,
    "official_api_documentation": 'https://github.com/thedaviddelta/lingva-translate#public-apis',
    "use_official_api": True,
    "require_api_key": False,
    "results": 'JSON',
}

engine_type = 'online_dictionary'
categories = ['general']

url = "https://lingva.ml"
search_url = "{url}/api/v1/{from_lang}/{to_lang}/{query}"


def request(_query, params):
    params['url'] = search_url.format(
        url=url, from_lang=params['from_lang'][1], to_lang=params['to_lang'][1], query=params['query']
    )
    return params


def response(resp):
    results = []

    result = loads(resp.text)
    info = result["info"]
    from_to_prefix = "%s-%s " % (resp.search_params['from_lang'][1], resp.search_params['to_lang'][1])

    if "typo" in info:
        results.append({"suggestion": from_to_prefix + info["typo"]})

    if 'definitions' in info:  # pylint: disable=too-many-nested-blocks
        for definition in info['definitions']:
            if 'list' in definition:
                for item in definition['list']:
                    if 'synonyms' in item:
                        for synonym in item['synonyms']:
                            results.append({"suggestion": from_to_prefix + synonym})

    infobox = ""

    for translation in info["extraTranslations"]:
        infobox += f"<b>{translation['type']}</b>"

        for word in translation["list"]:
            infobox += f"<dl><dt>{word['word']}</dt>"

            for meaning in word["meanings"]:
                infobox += f"<dd>{meaning}</dd>"

            infobox += "</dl>"

    results.append(
        {
            'infobox': result["translation"],
            'content': infobox,
        }
    )

    return results
pick engine fixes (#3306) * [fix] google engine: results XPath * [fix] google & youtube - set EU consent cookie This change the previous bypass method for Google consent using ``ucbcb=1`` (6face215b8) to accept the consent using ``CONSENT=YES+``. The youtube_noapi and google have a similar API, at least for the consent[1]. Get CONSENT cookie from google reguest:: curl -i "https://www.google.com/search?q=time&tbm=isch" \ -A "Mozilla/5.0 (X11; Linux i686; rv:102.0) Gecko/20100101 Firefox/102.0" \ \| grep -i consent ... location: https://consent.google.com/m?continue=https://www.google.com/search?q%3Dtime%26tbm%3Disch&gl=DE&m=0&pc=irp&uxe=eomtm&hl=en-US&src=1 set-cookie: CONSENT=PENDING+936; expires=Wed, 24-Jul-2024 11:26:20 GMT; path=/; domain=.google.com; Secure ... PENDING & YES [2]: Google change the way for consent about YouTube cookies agreement in EU countries. Instead of showing a popup in the website, YouTube redirects the user to a new webpage at consent.youtube.com domain ... Fix for this is to put a cookie CONSENT with YES+ value for every YouTube request [1] https://github.com/iv-org/invidious/pull/2207 [2] https://github.com/TeamNewPipe/NewPipeExtractor/issues/592 Closes: https://github.com/searxng/searxng/issues/1432 * [fix] sjp engine - convert enginename to a latin1 compliance name The engine name is not only a name its also a identifier that is used in logs, HTTP headers and more. Unicode characters in the name of an engine could cause various issues. Closes: https://github.com/searxng/searxng/issues/1544 Signed-off-by: Markus Heiser <markus.heiser@darmarit.de> * [fix] engine tineye: handle 422 response of not supported img format Closes: https://github.com/searxng/searxng/issues/1449 Signed-off-by: Markus Heiser <markus.heiser@darmarit.de> * bypass google consent with ucbcb=1 * [mod] Adds Lingva translate engine Add the lingva engine (which grabs data from google translate). Results from Lingva are added to the infobox results. * openstreetmap engine: return the localized named. For example: display "Tokyo" instead of "東京都" when the language is English. * [fix] engines/openstreetmap.py typo: user_langage --> user_language Signed-off-by: Markus Heiser <markus.heiser@darmarit.de> * Wikidata engine: ignore dummy entities * Wikidata engine: minor change of the SPARQL request The engine can be slow especially when the query won't return any answer. See https://www.mediawiki.org/wiki/Wikidata_Query_Service/User_Manual/MWAPI#Find_articles_in_Wikipedia_speaking_about_cheese_and_see_which_Wikibase_items_they_correspond_to Co-authored-by: Léon Tiekötter <leon@tiekoetter.com> Co-authored-by: Emilien Devos <contact@emiliendevos.be> Co-authored-by: Markus Heiser <markus.heiser@darmarit.de> Co-authored-by: Emilien Devos <github@emiliendevos.be> Co-authored-by: ta <alt3753.7@gmail.com> Co-authored-by: Alexandre Flament <alex@al-f.net> 2022-07-30 21:45:07 +02:00			`# SPDX-License-Identifier: AGPL-3.0-or-later`
			`# lint: pylint`
			`"""Lingva (alternative Google Translate frontend)"""`

			`from json import loads`

			`about = {`
			`"website": 'https://lingva.ml',`
			`"wikidata_id": None,`
			`"official_api_documentation": 'https://github.com/thedaviddelta/lingva-translate#public-apis',`
			`"use_official_api": True,`
			`"require_api_key": False,`
			`"results": 'JSON',`
			`}`

			`engine_type = 'online_dictionary'`
			`categories = ['general']`

			`url = "https://lingva.ml"`
			`search_url = "{url}/api/v1/{from_lang}/{to_lang}/{query}"`


			`def request(_query, params):`
			`params['url'] = search_url.format(`
			`url=url, from_lang=params['from_lang'][1], to_lang=params['to_lang'][1], query=params['query']`
			`)`
			`return params`


			`def response(resp):`
			`results = []`

			`result = loads(resp.text)`
			`info = result["info"]`
			`from_to_prefix = "%s-%s " % (resp.search_params['from_lang'][1], resp.search_params['to_lang'][1])`

			`if "typo" in info:`
			`results.append({"suggestion": from_to_prefix + info["typo"]})`

			`if 'definitions' in info: # pylint: disable=too-many-nested-blocks`
			`for definition in info['definitions']:`
			`if 'list' in definition:`
			`for item in definition['list']:`
			`if 'synonyms' in item:`
			`for synonym in item['synonyms']:`
			`results.append({"suggestion": from_to_prefix + synonym})`

			`infobox = ""`

			`for translation in info["extraTranslations"]:`
			`infobox += f"<b>{translation['type']}</b>"`

			`for word in translation["list"]:`
			`infobox += f"<dl><dt>{word['word']}</dt>"`

			`for meaning in word["meanings"]:`
			`infobox += f"<dd>{meaning}</dd>"`

			`infobox += "</dl>"`

			`results.append(`
			`{`
			`'infobox': result["translation"],`
			`'content': infobox,`
			`}`
			`)`

			`return results`