You can not select more than 25 topics Topics must start with a letter or number, can include dashes ('-') and can be up to 35 characters long.

595 lines
18KB

  1. /*
  2. * FLAC (Free Lossless Audio Codec) decoder
  3. * Copyright (c) 2003 Alex Beregszaszi
  4. *
  5. * This file is part of FFmpeg.
  6. *
  7. * FFmpeg is free software; you can redistribute it and/or
  8. * modify it under the terms of the GNU Lesser General Public
  9. * License as published by the Free Software Foundation; either
  10. * version 2.1 of the License, or (at your option) any later version.
  11. *
  12. * FFmpeg is distributed in the hope that it will be useful,
  13. * but WITHOUT ANY WARRANTY; without even the implied warranty of
  14. * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU
  15. * Lesser General Public License for more details.
  16. *
  17. * You should have received a copy of the GNU Lesser General Public
  18. * License along with FFmpeg; if not, write to the Free Software
  19. * Foundation, Inc., 51 Franklin Street, Fifth Floor, Boston, MA 02110-1301 USA
  20. */
  21. /**
  22. * @file
  23. * FLAC (Free Lossless Audio Codec) decoder
  24. * @author Alex Beregszaszi
  25. * @see http://flac.sourceforge.net/
  26. *
  27. * This decoder can be used in 1 of 2 ways: Either raw FLAC data can be fed
  28. * through, starting from the initial 'fLaC' signature; or by passing the
  29. * 34-byte streaminfo structure through avctx->extradata[_size] followed
  30. * by data starting with the 0xFFF8 marker.
  31. */
  32. #include <limits.h>
  33. #include "libavutil/avassert.h"
  34. #include "libavutil/channel_layout.h"
  35. #include "libavutil/crc.h"
  36. #include "avcodec.h"
  37. #include "internal.h"
  38. #include "get_bits.h"
  39. #include "bytestream.h"
  40. #include "golomb.h"
  41. #include "flac.h"
  42. #include "flacdata.h"
  43. #include "flacdsp.h"
  44. typedef struct FLACContext {
  45. FLACSTREAMINFO
  46. AVCodecContext *avctx; ///< parent AVCodecContext
  47. GetBitContext gb; ///< GetBitContext initialized to start at the current frame
  48. int blocksize; ///< number of samples in the current frame
  49. int sample_shift; ///< shift required to make output samples 16-bit or 32-bit
  50. int ch_mode; ///< channel decorrelation type in the current frame
  51. int got_streaminfo; ///< indicates if the STREAMINFO has been read
  52. int32_t *decoded[FLAC_MAX_CHANNELS]; ///< decoded samples
  53. uint8_t *decoded_buffer;
  54. unsigned int decoded_buffer_size;
  55. FLACDSPContext dsp;
  56. } FLACContext;
  57. static int allocate_buffers(FLACContext *s);
  58. static void flac_set_bps(FLACContext *s)
  59. {
  60. enum AVSampleFormat req = s->avctx->request_sample_fmt;
  61. int need32 = s->bps > 16;
  62. int want32 = av_get_bytes_per_sample(req) > 2;
  63. int planar = av_sample_fmt_is_planar(req);
  64. if (need32 || want32) {
  65. if (planar)
  66. s->avctx->sample_fmt = AV_SAMPLE_FMT_S32P;
  67. else
  68. s->avctx->sample_fmt = AV_SAMPLE_FMT_S32;
  69. s->sample_shift = 32 - s->bps;
  70. } else {
  71. if (planar)
  72. s->avctx->sample_fmt = AV_SAMPLE_FMT_S16P;
  73. else
  74. s->avctx->sample_fmt = AV_SAMPLE_FMT_S16;
  75. s->sample_shift = 16 - s->bps;
  76. }
  77. }
  78. static av_cold int flac_decode_init(AVCodecContext *avctx)
  79. {
  80. enum FLACExtradataFormat format;
  81. uint8_t *streaminfo;
  82. int ret;
  83. FLACContext *s = avctx->priv_data;
  84. s->avctx = avctx;
  85. /* for now, the raw FLAC header is allowed to be passed to the decoder as
  86. frame data instead of extradata. */
  87. if (!avctx->extradata)
  88. return 0;
  89. if (!avpriv_flac_is_extradata_valid(avctx, &format, &streaminfo))
  90. return -1;
  91. /* initialize based on the demuxer-supplied streamdata header */
  92. avpriv_flac_parse_streaminfo(avctx, (FLACStreaminfo *)s, streaminfo);
  93. ret = allocate_buffers(s);
  94. if (ret < 0)
  95. return ret;
  96. flac_set_bps(s);
  97. ff_flacdsp_init(&s->dsp, avctx->sample_fmt, s->bps);
  98. s->got_streaminfo = 1;
  99. return 0;
  100. }
  101. static void dump_headers(AVCodecContext *avctx, FLACStreaminfo *s)
  102. {
  103. av_log(avctx, AV_LOG_DEBUG, " Max Blocksize: %d\n", s->max_blocksize);
  104. av_log(avctx, AV_LOG_DEBUG, " Max Framesize: %d\n", s->max_framesize);
  105. av_log(avctx, AV_LOG_DEBUG, " Samplerate: %d\n", s->samplerate);
  106. av_log(avctx, AV_LOG_DEBUG, " Channels: %d\n", s->channels);
  107. av_log(avctx, AV_LOG_DEBUG, " Bits: %d\n", s->bps);
  108. }
  109. static int allocate_buffers(FLACContext *s)
  110. {
  111. int buf_size;
  112. av_assert0(s->max_blocksize);
  113. buf_size = av_samples_get_buffer_size(NULL, s->channels, s->max_blocksize,
  114. AV_SAMPLE_FMT_S32P, 0);
  115. if (buf_size < 0)
  116. return buf_size;
  117. av_fast_malloc(&s->decoded_buffer, &s->decoded_buffer_size, buf_size);
  118. if (!s->decoded_buffer)
  119. return AVERROR(ENOMEM);
  120. return av_samples_fill_arrays((uint8_t **)s->decoded, NULL,
  121. s->decoded_buffer, s->channels,
  122. s->max_blocksize, AV_SAMPLE_FMT_S32P, 0);
  123. }
  124. /**
  125. * Parse the STREAMINFO from an inline header.
  126. * @param s the flac decoding context
  127. * @param buf input buffer, starting with the "fLaC" marker
  128. * @param buf_size buffer size
  129. * @return non-zero if metadata is invalid
  130. */
  131. static int parse_streaminfo(FLACContext *s, const uint8_t *buf, int buf_size)
  132. {
  133. int metadata_type, metadata_size, ret;
  134. if (buf_size < FLAC_STREAMINFO_SIZE+8) {
  135. /* need more data */
  136. return 0;
  137. }
  138. avpriv_flac_parse_block_header(&buf[4], NULL, &metadata_type, &metadata_size);
  139. if (metadata_type != FLAC_METADATA_TYPE_STREAMINFO ||
  140. metadata_size != FLAC_STREAMINFO_SIZE) {
  141. return AVERROR_INVALIDDATA;
  142. }
  143. avpriv_flac_parse_streaminfo(s->avctx, (FLACStreaminfo *)s, &buf[8]);
  144. ret = allocate_buffers(s);
  145. if (ret < 0)
  146. return ret;
  147. flac_set_bps(s);
  148. ff_flacdsp_init(&s->dsp, s->avctx->sample_fmt, s->bps);
  149. s->got_streaminfo = 1;
  150. return 0;
  151. }
  152. /**
  153. * Determine the size of an inline header.
  154. * @param buf input buffer, starting with the "fLaC" marker
  155. * @param buf_size buffer size
  156. * @return number of bytes in the header, or 0 if more data is needed
  157. */
  158. static int get_metadata_size(const uint8_t *buf, int buf_size)
  159. {
  160. int metadata_last, metadata_size;
  161. const uint8_t *buf_end = buf + buf_size;
  162. buf += 4;
  163. do {
  164. if (buf_end - buf < 4)
  165. return 0;
  166. avpriv_flac_parse_block_header(buf, &metadata_last, NULL, &metadata_size);
  167. buf += 4;
  168. if (buf_end - buf < metadata_size) {
  169. /* need more data in order to read the complete header */
  170. return 0;
  171. }
  172. buf += metadata_size;
  173. } while (!metadata_last);
  174. return buf_size - (buf_end - buf);
  175. }
  176. static int decode_residuals(FLACContext *s, int32_t *decoded, int pred_order)
  177. {
  178. int i, tmp, partition, method_type, rice_order;
  179. int rice_bits, rice_esc;
  180. int samples;
  181. method_type = get_bits(&s->gb, 2);
  182. if (method_type > 1) {
  183. av_log(s->avctx, AV_LOG_ERROR, "illegal residual coding method %d\n",
  184. method_type);
  185. return -1;
  186. }
  187. rice_order = get_bits(&s->gb, 4);
  188. samples= s->blocksize >> rice_order;
  189. if (pred_order > samples) {
  190. av_log(s->avctx, AV_LOG_ERROR, "invalid predictor order: %i > %i\n",
  191. pred_order, samples);
  192. return -1;
  193. }
  194. rice_bits = 4 + method_type;
  195. rice_esc = (1 << rice_bits) - 1;
  196. decoded += pred_order;
  197. i= pred_order;
  198. for (partition = 0; partition < (1 << rice_order); partition++) {
  199. tmp = get_bits(&s->gb, rice_bits);
  200. if (tmp == rice_esc) {
  201. tmp = get_bits(&s->gb, 5);
  202. for (; i < samples; i++)
  203. *decoded++ = get_sbits_long(&s->gb, tmp);
  204. } else {
  205. for (; i < samples; i++) {
  206. *decoded++ = get_sr_golomb_flac(&s->gb, tmp, INT_MAX, 0);
  207. }
  208. }
  209. i= 0;
  210. }
  211. return 0;
  212. }
  213. static int decode_subframe_fixed(FLACContext *s, int32_t *decoded,
  214. int pred_order, int bps)
  215. {
  216. const int blocksize = s->blocksize;
  217. int av_uninit(a), av_uninit(b), av_uninit(c), av_uninit(d), i;
  218. /* warm up samples */
  219. for (i = 0; i < pred_order; i++) {
  220. decoded[i] = get_sbits_long(&s->gb, bps);
  221. }
  222. if (decode_residuals(s, decoded, pred_order) < 0)
  223. return -1;
  224. if (pred_order > 0)
  225. a = decoded[pred_order-1];
  226. if (pred_order > 1)
  227. b = a - decoded[pred_order-2];
  228. if (pred_order > 2)
  229. c = b - decoded[pred_order-2] + decoded[pred_order-3];
  230. if (pred_order > 3)
  231. d = c - decoded[pred_order-2] + 2*decoded[pred_order-3] - decoded[pred_order-4];
  232. switch (pred_order) {
  233. case 0:
  234. break;
  235. case 1:
  236. for (i = pred_order; i < blocksize; i++)
  237. decoded[i] = a += decoded[i];
  238. break;
  239. case 2:
  240. for (i = pred_order; i < blocksize; i++)
  241. decoded[i] = a += b += decoded[i];
  242. break;
  243. case 3:
  244. for (i = pred_order; i < blocksize; i++)
  245. decoded[i] = a += b += c += decoded[i];
  246. break;
  247. case 4:
  248. for (i = pred_order; i < blocksize; i++)
  249. decoded[i] = a += b += c += d += decoded[i];
  250. break;
  251. default:
  252. av_log(s->avctx, AV_LOG_ERROR, "illegal pred order %d\n", pred_order);
  253. return -1;
  254. }
  255. return 0;
  256. }
  257. static int decode_subframe_lpc(FLACContext *s, int32_t *decoded, int pred_order,
  258. int bps)
  259. {
  260. int i;
  261. int coeff_prec, qlevel;
  262. int coeffs[32];
  263. /* warm up samples */
  264. for (i = 0; i < pred_order; i++) {
  265. decoded[i] = get_sbits_long(&s->gb, bps);
  266. }
  267. coeff_prec = get_bits(&s->gb, 4) + 1;
  268. if (coeff_prec == 16) {
  269. av_log(s->avctx, AV_LOG_ERROR, "invalid coeff precision\n");
  270. return -1;
  271. }
  272. qlevel = get_sbits(&s->gb, 5);
  273. if (qlevel < 0) {
  274. av_log(s->avctx, AV_LOG_ERROR, "qlevel %d not supported, maybe buggy stream\n",
  275. qlevel);
  276. return -1;
  277. }
  278. for (i = 0; i < pred_order; i++) {
  279. coeffs[pred_order - i - 1] = get_sbits(&s->gb, coeff_prec);
  280. }
  281. if (decode_residuals(s, decoded, pred_order) < 0)
  282. return -1;
  283. s->dsp.lpc(decoded, coeffs, pred_order, qlevel, s->blocksize);
  284. return 0;
  285. }
  286. static inline int decode_subframe(FLACContext *s, int channel)
  287. {
  288. int32_t *decoded = s->decoded[channel];
  289. int type, wasted = 0;
  290. int bps = s->bps;
  291. int i, tmp;
  292. if (channel == 0) {
  293. if (s->ch_mode == FLAC_CHMODE_RIGHT_SIDE)
  294. bps++;
  295. } else {
  296. if (s->ch_mode == FLAC_CHMODE_LEFT_SIDE || s->ch_mode == FLAC_CHMODE_MID_SIDE)
  297. bps++;
  298. }
  299. if (get_bits1(&s->gb)) {
  300. av_log(s->avctx, AV_LOG_ERROR, "invalid subframe padding\n");
  301. return -1;
  302. }
  303. type = get_bits(&s->gb, 6);
  304. if (get_bits1(&s->gb)) {
  305. int left = get_bits_left(&s->gb);
  306. wasted = 1;
  307. if ( left < 0 ||
  308. (left < bps && !show_bits_long(&s->gb, left)) ||
  309. !show_bits_long(&s->gb, bps)) {
  310. av_log(s->avctx, AV_LOG_ERROR,
  311. "Invalid number of wasted bits > available bits (%d) - left=%d\n",
  312. bps, left);
  313. return AVERROR_INVALIDDATA;
  314. }
  315. while (!get_bits1(&s->gb))
  316. wasted++;
  317. bps -= wasted;
  318. }
  319. if (bps > 32) {
  320. av_log_missing_feature(s->avctx, "Decorrelated bit depth > 32", 0);
  321. return AVERROR_PATCHWELCOME;
  322. }
  323. //FIXME use av_log2 for types
  324. if (type == 0) {
  325. tmp = get_sbits_long(&s->gb, bps);
  326. for (i = 0; i < s->blocksize; i++)
  327. decoded[i] = tmp;
  328. } else if (type == 1) {
  329. for (i = 0; i < s->blocksize; i++)
  330. decoded[i] = get_sbits_long(&s->gb, bps);
  331. } else if ((type >= 8) && (type <= 12)) {
  332. if (decode_subframe_fixed(s, decoded, type & ~0x8, bps) < 0)
  333. return -1;
  334. } else if (type >= 32) {
  335. if (decode_subframe_lpc(s, decoded, (type & ~0x20)+1, bps) < 0)
  336. return -1;
  337. } else {
  338. av_log(s->avctx, AV_LOG_ERROR, "invalid coding type\n");
  339. return -1;
  340. }
  341. if (wasted) {
  342. int i;
  343. for (i = 0; i < s->blocksize; i++)
  344. decoded[i] <<= wasted;
  345. }
  346. return 0;
  347. }
  348. static int decode_frame(FLACContext *s)
  349. {
  350. int i, ret;
  351. GetBitContext *gb = &s->gb;
  352. FLACFrameInfo fi;
  353. if (ff_flac_decode_frame_header(s->avctx, gb, &fi, 0)) {
  354. av_log(s->avctx, AV_LOG_ERROR, "invalid frame header\n");
  355. return -1;
  356. }
  357. if (s->channels && fi.channels != s->channels && s->got_streaminfo) {
  358. s->channels = s->avctx->channels = fi.channels;
  359. ff_flac_set_channel_layout(s->avctx);
  360. ret = allocate_buffers(s);
  361. if (ret < 0)
  362. return ret;
  363. }
  364. s->channels = s->avctx->channels = fi.channels;
  365. if (!s->avctx->channel_layout)
  366. ff_flac_set_channel_layout(s->avctx);
  367. s->ch_mode = fi.ch_mode;
  368. if (!s->bps && !fi.bps) {
  369. av_log(s->avctx, AV_LOG_ERROR, "bps not found in STREAMINFO or frame header\n");
  370. return -1;
  371. }
  372. if (!fi.bps) {
  373. fi.bps = s->bps;
  374. } else if (s->bps && fi.bps != s->bps) {
  375. av_log(s->avctx, AV_LOG_ERROR, "switching bps mid-stream is not "
  376. "supported\n");
  377. return -1;
  378. }
  379. if (!s->bps) {
  380. s->bps = s->avctx->bits_per_raw_sample = fi.bps;
  381. flac_set_bps(s);
  382. }
  383. if (!s->max_blocksize)
  384. s->max_blocksize = FLAC_MAX_BLOCKSIZE;
  385. if (fi.blocksize > s->max_blocksize) {
  386. av_log(s->avctx, AV_LOG_ERROR, "blocksize %d > %d\n", fi.blocksize,
  387. s->max_blocksize);
  388. return -1;
  389. }
  390. s->blocksize = fi.blocksize;
  391. if (!s->samplerate && !fi.samplerate) {
  392. av_log(s->avctx, AV_LOG_ERROR, "sample rate not found in STREAMINFO"
  393. " or frame header\n");
  394. return -1;
  395. }
  396. if (fi.samplerate == 0)
  397. fi.samplerate = s->samplerate;
  398. s->samplerate = s->avctx->sample_rate = fi.samplerate;
  399. if (!s->got_streaminfo) {
  400. ret = allocate_buffers(s);
  401. if (ret < 0)
  402. return ret;
  403. ff_flacdsp_init(&s->dsp, s->avctx->sample_fmt, s->bps);
  404. s->got_streaminfo = 1;
  405. dump_headers(s->avctx, (FLACStreaminfo *)s);
  406. }
  407. // dump_headers(s->avctx, (FLACStreaminfo *)s);
  408. /* subframes */
  409. for (i = 0; i < s->channels; i++) {
  410. if (decode_subframe(s, i) < 0)
  411. return -1;
  412. }
  413. align_get_bits(gb);
  414. /* frame footer */
  415. skip_bits(gb, 16); /* data crc */
  416. return 0;
  417. }
  418. static int flac_decode_frame(AVCodecContext *avctx, void *data,
  419. int *got_frame_ptr, AVPacket *avpkt)
  420. {
  421. AVFrame *frame = data;
  422. const uint8_t *buf = avpkt->data;
  423. int buf_size = avpkt->size;
  424. FLACContext *s = avctx->priv_data;
  425. int bytes_read = 0;
  426. int ret;
  427. *got_frame_ptr = 0;
  428. if (s->max_framesize == 0) {
  429. s->max_framesize =
  430. ff_flac_get_max_frame_size(s->max_blocksize ? s->max_blocksize : FLAC_MAX_BLOCKSIZE,
  431. FLAC_MAX_CHANNELS, 32);
  432. }
  433. if (buf_size > 5 && !memcmp(buf, "\177FLAC", 5)) {
  434. av_log(s->avctx, AV_LOG_DEBUG, "skiping flac header packet 1\n");
  435. return buf_size;
  436. }
  437. if (buf_size > 0 && (*buf & 0x7F) == FLAC_METADATA_TYPE_VORBIS_COMMENT) {
  438. av_log(s->avctx, AV_LOG_DEBUG, "skiping vorbis comment\n");
  439. return buf_size;
  440. }
  441. /* check that there is at least the smallest decodable amount of data.
  442. this amount corresponds to the smallest valid FLAC frame possible.
  443. FF F8 69 02 00 00 9A 00 00 34 46 */
  444. if (buf_size < FLAC_MIN_FRAME_SIZE)
  445. return buf_size;
  446. /* check for inline header */
  447. if (AV_RB32(buf) == MKBETAG('f','L','a','C')) {
  448. if (!s->got_streaminfo && parse_streaminfo(s, buf, buf_size)) {
  449. av_log(s->avctx, AV_LOG_ERROR, "invalid header\n");
  450. return -1;
  451. }
  452. return get_metadata_size(buf, buf_size);
  453. }
  454. /* decode frame */
  455. init_get_bits(&s->gb, buf, buf_size*8);
  456. if (decode_frame(s) < 0) {
  457. av_log(s->avctx, AV_LOG_ERROR, "decode_frame() failed\n");
  458. return -1;
  459. }
  460. bytes_read = get_bits_count(&s->gb)/8;
  461. if ((s->avctx->err_recognition & AV_EF_CRCCHECK) &&
  462. av_crc(av_crc_get_table(AV_CRC_16_ANSI),
  463. 0, buf, bytes_read)) {
  464. av_log(s->avctx, AV_LOG_ERROR, "CRC error at PTS %"PRId64"\n", avpkt->pts);
  465. if (s->avctx->err_recognition & AV_EF_EXPLODE)
  466. return AVERROR_INVALIDDATA;
  467. }
  468. /* get output buffer */
  469. frame->nb_samples = s->blocksize;
  470. if ((ret = ff_get_buffer(avctx, frame, 0)) < 0) {
  471. av_log(avctx, AV_LOG_ERROR, "get_buffer() failed\n");
  472. return ret;
  473. }
  474. s->dsp.decorrelate[s->ch_mode](frame->data, s->decoded, s->channels,
  475. s->blocksize, s->sample_shift);
  476. if (bytes_read > buf_size) {
  477. av_log(s->avctx, AV_LOG_ERROR, "overread: %d\n", bytes_read - buf_size);
  478. return -1;
  479. }
  480. if (bytes_read < buf_size) {
  481. av_log(s->avctx, AV_LOG_DEBUG, "underread: %d orig size: %d\n",
  482. buf_size - bytes_read, buf_size);
  483. }
  484. *got_frame_ptr = 1;
  485. return bytes_read;
  486. }
  487. static av_cold int flac_decode_close(AVCodecContext *avctx)
  488. {
  489. FLACContext *s = avctx->priv_data;
  490. av_freep(&s->decoded_buffer);
  491. return 0;
  492. }
  493. AVCodec ff_flac_decoder = {
  494. .name = "flac",
  495. .type = AVMEDIA_TYPE_AUDIO,
  496. .id = AV_CODEC_ID_FLAC,
  497. .priv_data_size = sizeof(FLACContext),
  498. .init = flac_decode_init,
  499. .close = flac_decode_close,
  500. .decode = flac_decode_frame,
  501. .capabilities = CODEC_CAP_DR1,
  502. .long_name = NULL_IF_CONFIG_SMALL("FLAC (Free Lossless Audio Codec)"),
  503. .sample_fmts = (const enum AVSampleFormat[]) { AV_SAMPLE_FMT_S16,
  504. AV_SAMPLE_FMT_S16P,
  505. AV_SAMPLE_FMT_S32,
  506. AV_SAMPLE_FMT_S32P,
  507. AV_SAMPLE_FMT_NONE },
  508. };