[ard] Add suppor for plain ARD downloads (Fixes #3546)

author Philipp Hagemeister <phihag@phihag.de>

Wed, 27 Aug 2014 00:36:57 +0000 (02:36 +0200)

committer Philipp Hagemeister <phihag@phihag.de>

Wed, 27 Aug 2014 00:36:57 +0000 (02:36 +0200)
author Philipp Hagemeister <phihag@phihag.de>
Wed, 27 Aug 2014 00:36:57 +0000 (02:36 +0200)
committer Philipp Hagemeister <phihag@phihag.de>
Wed, 27 Aug 2014 00:36:57 +0000 (02:36 +0200)
diff --git a/youtube_dl/extractor/__init__.py b/youtube_dl/extractor/__init__.py

index d26454396785bf8adc04e1b72cbd3a42260cc204..17ed69ba4205a4aa91fc989c7d7b5f02dbc482ed 100644 (file)
--- a/youtube_dl/extractor/__init__.py
+++ b/youtube_dl/extractor/__init__.py
@@ -9,7 +9,7 @@ from .allocine import AllocineIE
  from .aparat import AparatIE
  from .appletrailers import AppleTrailersIE
  from .archiveorg import ArchiveOrgIE
-from .ard import ARDIE
+from .ard import ARDIE, ARDMediathekIE
  from .arte import (
      ArteTvIE,
      ArteTVPlus7IE,
diff --git a/youtube_dl/extractor/ard.py b/youtube_dl/extractor/ard.py

index 7f0da8ab6d5b9f0e62f2af18c74c26185c505259..ef94c72395723b31bd444e80b6ba12d990acf38b 100644 (file)
--- a/youtube_dl/extractor/ard.py
+++ b/youtube_dl/extractor/ard.py
@@ -10,10 +10,14 @@ from ..utils import (
      qualities,
      compat_urllib_parse_urlparse,
      compat_urllib_parse,
+    int_or_none,
+    parse_duration,
+    unified_strdate,
  )
  
  
-class ARDIE(InfoExtractor):
+class ARDMediathekIE(InfoExtractor):
+    IE_NAME = 'ARD:mediathek'
      _VALID_URL = r'^https?://(?:(?:www\.)?ardmediathek\.de|mediathek\.daserste\.de)/(?:.*/)(?P<video_id>[0-9]+|[^0-9][^/\?]+)[^/\?]*(?:\?.*)?'
  
      _TESTS = [{
@@ -128,3 +132,60 @@ class ARDIE(InfoExtractor):
              'formats': formats,
              'thumbnail': thumbnail,
          }
+
+
+class ARDIE(InfoExtractor):
+    _VALID_URL = '(?P<mainurl>https?://(www\.)?daserste\.de/[^?#]+/videos/(?P<display_id>[^/?#]+)-(?P<id>[0-9]+))\.html'
+    _TEST = {
+        'url': 'http://www.daserste.de/information/reportage-dokumentation/dokus/videos/die-story-im-ersten-mission-unter-falscher-flagge-100.html',
+        'md5': 'd216c3a86493f9322545e045ddc3eb35',
+        'info_dict': {
+            'display_id': 'die-story-im-ersten-mission-unter-falscher-flagge',
+            'id': '100',
+            'ext': 'mp4',
+            'duration': 2600,
+            'title': 'Die Story im Ersten: Mission unter falscher Flagge',
+            'upload_date': '20140804',
+            'thumbnail': 're:^https?://.*\.jpg$',
+        }
+    }
+
+    def _real_extract(self, url):
+        mobj = re.match(self._VALID_URL, url)
+        display_id = mobj.group('display_id')
+
+        player_url = mobj.group('mainurl') + '~playerXml.xml'
+        doc = self._download_xml(player_url, display_id)
+        video_node = doc.find('./video')
+        upload_date = unified_strdate(video_node.find('./broadcastDate').text)
+        thumbnail = video_node.find('.//teaserImage//variant/url').text
+
+        formats = []
+        for a in video_node.findall('.//asset'):
+            f = {
+                'format_id': a.attrib['type'],
+                'width': int_or_none(a.find('./frameWidth').text),
+                'height': int_or_none(a.find('./frameHeight').text),
+                'vbr': int_or_none(a.find('./bitrateVideo').text),
+                'abr': int_or_none(a.find('./bitrateAudio').text),
+                'vcodec': a.find('./codecVideo').text,
+                'tbr': int_or_none(a.find('./totalBitrate').text),
+            }
+            if a.find('./serverPrefix').text:
+                f['url'] = a.find('./serverPrefix').text
+                f['playpath'] = a.find('./fileName').text
+            else:
+                f['url'] = a.find('./fileName').text
+            formats.append(f)
+        self._sort_formats(formats)
+
+        return {
+            'id': mobj.group('id'),
+            'formats': formats,
+            'display_id': display_id,
+            'title': video_node.find('./title').text,
+            'duration': parse_duration(video_node.find('./duration').text),
+            'upload_date': upload_date,
+            'thumbnail': thumbnail,
+        }
+
author	Philipp Hagemeister <phihag@phihag.de>
	Wed, 27 Aug 2014 00:36:57 +0000 (02:36 +0200)
committer	Philipp Hagemeister <phihag@phihag.de>
	Wed, 27 Aug 2014 00:36:57 +0000 (02:36 +0200)
youtube_dl/extractor/__init__.py		patch \| blob \| history
youtube_dl/extractor/ard.py		patch \| blob \| history