Add support for https for all extractors as preventive and future-proof measure

[youtube-dl.git] / youtube_dl / extractor / arte.py
diff --git a/youtube_dl/extractor/arte.py b/youtube_dl/extractor/arte.py

index efde7e207bc8d166e80f2a26429797684535d114..ae0f27dcbe059c0d469eaeca243ef59400ff68d6 100644 (file)
--- a/youtube_dl/extractor/arte.py
+++ b/youtube_dl/extractor/arte.py
@@ -23,7 +23,7 @@ from ..utils import (
  
  
  class ArteTvIE(InfoExtractor):
-    _VALID_URL = r'http://videos\.arte\.tv/(?P<lang>fr|de|en|es)/.*-(?P<id>.*?)\.html'
+    _VALID_URL = r'https?://videos\.arte\.tv/(?P<lang>fr|de|en|es)/.*-(?P<id>.*?)\.html'
      IE_NAME = 'arte.tv'
  
      def _real_extract(self, url):
@@ -121,15 +121,18 @@ class ArteTVPlus7IE(InfoExtractor):
                  json_url = compat_parse_qs(
                      compat_urllib_parse_urlparse(iframe_url).query)['json_url'][0]
          if json_url:
-            return self._extract_from_json_url(json_url, video_id, lang)
-        # Differend kind of embed URL (e.g.
+            title = self._search_regex(
+                r'<h3[^>]+title=(["\'])(?P<title>.+?)\1',
+                webpage, 'title', default=None, group='title')
+            return self._extract_from_json_url(json_url, video_id, lang, title=title)
+        # Different kind of embed URL (e.g.
          # http://www.arte.tv/magazine/trepalium/fr/episode-0406-replay-trepalium)
          embed_url = self._search_regex(
              r'<iframe[^>]+src=(["\'])(?P<url>.+?)\1',
              webpage, 'embed url', group='url')
          return self.url_result(embed_url)
  
-    def _extract_from_json_url(self, json_url, video_id, lang):
+    def _extract_from_json_url(self, json_url, video_id, lang, title=None):
          info = self._download_json(json_url, video_id)
          player_info = info['videoJsonPlayer']
  
@@ -137,7 +140,7 @@ class ArteTVPlus7IE(InfoExtractor):
          if not upload_date_str:
              upload_date_str = (player_info.get('VRA') or player_info.get('VDA') or '').split(' ')[0]
  
-        title = player_info['VTI'].strip()
+        title = (player_info.get('VTI') or title or player_info['VID']).strip()
          subtitle = player_info.get('VSU', '').strip()
          if subtitle:
              title += ' - %s' % subtitle