Merge 13ed574168 into f9d98509a8

[BiliBiliBangumiIE] supports play_info extraction from webpage
- i.e. extracts premiums formats with logged-in cookies, haven't tested with format `12240` yet. * test url: https://www.bilibili.com/bangumi/play/ep829434, cookies: logged-in, non-premium
2024-11-25 00:31:26 +01:00 · 2024-11-17 23:45:53 +04:00 · 2024-08-20 00:46:21 +12:00 · 2024-08-17 10:58:16 +12:00 · 2024-07-27 23:10:53 +12:00 · 2024-07-27 22:51:15 +12:00
1 changed files with 25 additions and 14 deletions
--- a/yt_dlp/extractor/bilibili.py
+++ b/yt_dlp/extractor/bilibili.py
@ -165,14 +165,18 @@ class BilibiliBaseIE(InfoExtractor):
        params['w_rid'] = hashlib.md5(f'{query}{self._get_wbi_key(video_id)}'.encode()).hexdigest()
        return params
-    def _download_playinfo(self, bvid, cid, headers=None, qn=None):
+    def _download_playinfo(self, bvid, cid, headers=None, **kwargs):
-        params = {'bvid': bvid, 'cid': cid, 'fnval': 4048}
+        params = {'bvid': bvid, 'cid': cid, 'fnval': 4048, **kwargs}
-        if qn:
+        if self.is_logged_in:
-            params['qn'] = qn
+            params.pop('try_look', None)
        if kwargs.get('qn'):
            note = f'Downloading video format {kwargs["qn"]} for cid {cid}'
        else:
            note = f'Downloading video formats for cid {cid}'
        return self._download_json(
            'https://api.bilibili.com/x/player/wbi/playurl', bvid,
-            query=self._sign_wbi(params, bvid), headers=headers,
+            query=self._sign_wbi(params, bvid), headers=headers, note=note)['data']
            note=f'Downloading video formats for cid {cid} {qn or ""}')['data']
    def json2srt(self, json_data):
        srt_data = ''
@ -286,7 +290,7 @@ class BilibiliBaseIE(InfoExtractor):
            ('data', 'interaction', 'graph_version', {int_or_none}))
        cid_edges = self._get_divisions(video_id, graph_version, {1: {'cid': cid}}, 1)
        for cid, edges in cid_edges.items():
-            play_info = self._download_playinfo(video_id, cid, headers=headers)
+            play_info = self._download_playinfo(video_id, cid, headers=headers, try_look=1)
            yield {
                **metainfo,
                'id': f'{video_id}_{cid}',
@ -688,11 +692,12 @@ class BiliBiliIE(BilibiliBaseIE):
        aid = video_data.get('aid')
        old_video_id = format_field(aid, None, f'%s_part{part_id or 1}')
        cid = traverse_obj(video_data, ('pages', part_id - 1, 'cid')) if part_id else video_data.get('cid')
        if is_festival or not self.is_logged_in:
            query = {'try_look': 1} if not self.is_logged_in else {}
            play_info = self._download_playinfo(video_id, cid, headers=headers, **query)
        festival_info = {}
        if is_festival:
            play_info = self._download_playinfo(video_id, cid, headers=headers)
            festival_info = traverse_obj(initial_state, {
                'uploader': ('videoInfo', 'upName'),
                'uploader_id': ('videoInfo', 'upMid', {str_or_none}),
@ -730,7 +735,7 @@ class BiliBiliIE(BilibiliBaseIE):
        else:
            formats = self.extract_formats(play_info)
-            if not traverse_obj(play_info, ('dash')):
+            if not play_info.get('dash'):
                # we only have legacy formats and need additional work
                has_qn = lambda x: x in traverse_obj(formats, (..., 'quality'))
                for qn in traverse_obj(play_info, ('accept_quality', lambda _, v: not has_qn(v), {int})):
@ -860,10 +865,16 @@ class BiliBiliBangumiIE(BilibiliBaseIE):
            self.raise_login_required('This video is for premium members only')
        headers['Referer'] = url
        play_info = self._search_json(
            r'playurlSSRData\s*?=\s*?', webpage, 'embedded page info', episode_id,
            end_pattern='\n', default=None)
        if not play_info:
            play_info = self._download_json(
                'https://api.bilibili.com/pgc/player/web/v2/playurl', episode_id,
-            'Extracting episode', query={'fnval': '4048', 'ep_id': episode_id},
+                'Extracting episode', query={'fnval': 12240, 'ep_id': episode_id},
                headers=headers)
        premium_only = play_info.get('code') == -10403
        play_info = traverse_obj(play_info, ('result', 'video_info', {dict})) or {}
Author	SHA1	Message	Date
N/Ame	9286a4011b	Merge `13ed574168` into `f9d98509a8`	2024-11-17 23:45:53 +04:00
grqx_wsl	13ed574168	[BiliBiliBangumiIE] supports play_info extraction from webpage - i.e. extracts premiums formats with logged-in cookies, haven't tested with format `12240` yet. * test url: https://www.bilibili.com/bangumi/play/ep829434, cookies: logged-in, non-premium	2024-08-20 00:46:21 +12:00
grqx_wsl	79bb63957d	Merge remote-tracking branch 'upstream' into biliTryLook	2024-08-17 10:58:16 +12:00
grqx_wsl	d5dbdbccd3	`_download_playinfo`: more understandable note	2024-07-27 23:10:53 +12:00
grqx_wsl	b2965fa3b2	[BiliBiliBangumiIE] support format 12240(format name 智能修复, premium only) [cleanup]code formatting	2024-07-27 22:51:15 +12:00
grqx_wsl	510e29a42c	add support for _get_interactive_entries	2024-07-27 22:09:44 +12:00
grqx_wsl	90f4203632	keep the original `play_info` traversal	2024-07-26 10:46:41 +12:00
grqx_wsl	b01183f904	pops param `try_look` when logged in.	2024-07-26 10:04:18 +12:00
grqx_wsl	29a5968278	- Applied try_look to festival videos - Removed redundant calls to `_download_playinfo`	2024-07-26 03:07:32 +12:00
grqx_wsl	e187799c58	patch from https://github.com/yt-dlp/yt-dlp/issues/10554#issuecomment-2250014807 modified: yt_dlp/extractor/bilibili.py	2024-07-26 02:36:04 +12:00