yt-dlp/yt_dlp/extractor/godtube.py

from .common import InfoExtractor
from ..utils import (
    parse_duration,
    parse_iso8601,
)


class GodTubeIE(InfoExtractor):
    _WORKING = False
    _VALID_URL = r'https?://(?:www\.)?godtube\.com/watch/\?v=(?P<id>[\da-zA-Z]+)'
    _TESTS = [
        {
            'url': 'https://www.godtube.com/watch/?v=0C0CNNNU',
            'md5': '77108c1e4ab58f48031101a1a2119789',
            'info_dict': {
                'id': '0C0CNNNU',
                'ext': 'mp4',
                'title': 'Woman at the well.',
                'duration': 159,
                'timestamp': 1205712000,
                'uploader': 'beverlybmusic',
                'upload_date': '20080317',
                'thumbnail': r're:^https?://.*\.jpg$',
            },
        },
    ]

    def _real_extract(self, url):
        mobj = self._match_valid_url(url)
        video_id = mobj.group('id')

        config = self._download_xml(
            f'http://www.godtube.com/resource/mediaplayer/{video_id.lower()}.xml',
            video_id, 'Downloading player config XML')

        video_url = config.find('file').text
        uploader = config.find('author').text
        timestamp = parse_iso8601(config.find('date').text)
        duration = parse_duration(config.find('duration').text)
        thumbnail = config.find('image').text

        media = self._download_xml(
            f'http://www.godtube.com/media/xml/?v={video_id}', video_id, 'Downloading media XML')

        title = media.find('title').text

        return {
            'id': video_id,
            'url': video_url,
            'title': title,
            'thumbnail': thumbnail,
            'timestamp': timestamp,
            'uploader': uploader,
            'duration': duration,
        }
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00			`from .common import InfoExtractor`
			`from ..utils import (`
			`parse_duration,`
			`parse_iso8601,`
			`)`


			`class GodTubeIE(InfoExtractor):`
[cleanup] Mark broken and remove dead extractors (#9238) Authored by: seproDev 2024-03-08 19:02:45 -05:00			`_WORKING = False`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00			`_VALID_URL = r'https?://(?:www\.)?godtube\.com/watch/\?v=(?P<id>[\da-zA-Z]+)'`
			`_TESTS = [`
			`{`
			`'url': 'https://www.godtube.com/watch/?v=0C0CNNNU',`
			`'md5': '77108c1e4ab58f48031101a1a2119789',`
			`'info_dict': {`
			`'id': '0C0CNNNU',`
			`'ext': 'mp4',`
			`'title': 'Woman at the well.',`
			`'duration': 159,`
			`'timestamp': 1205712000,`
			`'uploader': 'beverlybmusic',`
			`'upload_date': '20080317',`
Fix "invalid escape sequences" error on Python 3.6 2017-01-02 07:08:07 -05:00			`'thumbnail': r're:^https?://.*\.jpg$',`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00			`},`
			`},`
			`]`

			`def _real_extract(self, url):`
[extractor] Common function `_match_valid_url` 2021-08-18 21:41:24 -04:00			`mobj = self._match_valid_url(url)`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00			`video_id = mobj.group('id')`

			`config = self._download_xml(`
[cleanup] Add more ruff rules (#10149) Authored by: seproDev Reviewed-by: bashonly <88596187+bashonly@users.noreply.github.com> Reviewed-by: Simon Sawicki <contact@grub4k.xyz> 2024-06-11 19:09:58 -04:00			`f'http://www.godtube.com/resource/mediaplayer/{video_id.lower()}.xml',`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00			`video_id, 'Downloading player config XML')`

[godtube] Fix on Python 2.6 2014-09-28 23:51:41 -04:00			`video_url = config.find('file').text`
			`uploader = config.find('author').text`
			`timestamp = parse_iso8601(config.find('date').text)`
			`duration = parse_duration(config.find('duration').text)`
			`thumbnail = config.find('image').text`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00
			`media = self._download_xml(`
[cleanup] Add more ruff rules (#10149) Authored by: seproDev Reviewed-by: bashonly <88596187+bashonly@users.noreply.github.com> Reviewed-by: Simon Sawicki <contact@grub4k.xyz> 2024-06-11 19:09:58 -04:00			`f'http://www.godtube.com/media/xml/?v={video_id}', video_id, 'Downloading media XML')`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00
[godtube] Fix on Python 2.6 2014-09-28 23:51:41 -04:00			`title = media.find('title').text`
[godtube] Add extractor (Closes #3367) 2014-07-26 15:38:05 -04:00
			`return {`
			`'id': video_id,`
			`'url': video_url,`
			`'title': title,`
			`'thumbnail': thumbnail,`
			`'timestamp': timestamp,`
			`'uploader': uploader,`
			`'duration': duration,`
			`}`