2 * Copyright 2012 Michael Chen <omxcodec@gmail.com>
3 * Copyright 2015 The CyanogenMod Project
5 * Licensed under the Apache License, Version 2.0 (the "License");
6 * you may not use this file except in compliance with the License.
7 * You may obtain a copy of the License at
9 * http://www.apache.org/licenses/LICENSE-2.0
11 * Unless required by applicable law or agreed to in writing, software
12 * distributed under the License is distributed on an "AS IS" BASIS,
13 * WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
14 * See the License for the specific language governing permissions and
15 * limitations under the License.
18 //#define LOG_NDEBUG 0
19 #define LOG_TAG "FFmpegExtractor"
20 #include <utils/Log.h>
23 #include <limits.h> /* INT_MAX */
25 #include <sys/prctl.h>
27 #include <utils/misc.h>
28 #include <utils/String8.h>
29 #include <cutils/properties.h>
30 #include <media/stagefright/foundation/ABitReader.h>
31 #include <media/stagefright/foundation/ABuffer.h>
32 #include <media/stagefright/foundation/ADebug.h>
33 #include <media/stagefright/foundation/AMessage.h>
34 #include <media/stagefright/foundation/hexdump.h>
35 #include <media/stagefright/DataSource.h>
36 #include <media/stagefright/MediaBuffer.h>
37 #include <media/stagefright/foundation/ADebug.h>
38 #include <media/stagefright/MediaDefs.h>
39 #include <media/stagefright/MediaErrors.h>
40 #include <media/stagefright/MediaSource.h>
41 #include <media/stagefright/MetaData.h>
42 #include <media/stagefright/Utils.h>
43 #include "include/avc_utils.h"
45 #include "utils/codec_utils.h"
46 #include "utils/ffmpeg_cmdutils.h"
48 #include "FFmpegExtractor.h"
50 #define MAX_QUEUE_SIZE (15 * 1024 * 1024)
51 #define MIN_AUDIOQ_SIZE (20 * 16 * 1024)
53 #define EXTRACTOR_MAX_PROBE_PACKETS 200
54 #define FF_MAX_EXTRADATA_SIZE ((1 << 28) - FF_INPUT_BUFFER_PADDING_SIZE)
56 #define WAIT_KEY_PACKET_AFTER_SEEK 1
57 #define SUPPOURT_UNKNOWN_FORMAT 1
60 #define DEBUG_READ_ENTRY 0
61 #define DEBUG_DISABLE_VIDEO 0
62 #define DEBUG_DISABLE_AUDIO 0
64 #define DEBUG_FORMATS 0
73 struct FFmpegSource : public MediaSource {
74 FFmpegSource(const sp<FFmpegExtractor> &extractor, size_t index);
76 virtual status_t start(MetaData *params);
77 virtual status_t stop();
78 virtual sp<MetaData> getFormat();
80 virtual status_t read(
81 MediaBuffer **buffer, const ReadOptions *options);
84 virtual ~FFmpegSource();
87 friend struct FFmpegExtractor;
89 sp<FFmpegExtractor> mExtractor;
92 enum AVMediaType mMediaType;
97 size_t mNALLengthSize;
103 int64_t mFirstKeyPktTimestamp;
107 DISALLOW_EVIL_CONSTRUCTORS(FFmpegSource);
110 ////////////////////////////////////////////////////////////////////////////////
112 FFmpegExtractor::FFmpegExtractor(const sp<DataSource> &source, const sp<AMessage> &meta)
113 : mDataSource(source),
116 mFFmpegInited(false),
118 mReaderThreadStarted(false) {
119 ALOGV("FFmpegExtractor::FFmpegExtractor");
121 fetchStuffsFromSniffedMeta(meta);
123 int err = initStreams();
125 ALOGE("failed to init ffmpeg");
129 // start reader here, as we want to extract extradata from bitstream if no extradata
132 while(mProbePkts <= EXTRACTOR_MAX_PROBE_PACKETS && !mEOF &&
133 (mFormatCtx->pb ? !mFormatCtx->pb->error : 1) &&
134 (mDefersToCreateVideoTrack || mDefersToCreateAudioTrack)) {
135 ALOGV("mProbePkts=%d", mProbePkts);
139 ALOGV("mProbePkts: %d, mEOF: %d, pb->error(if has): %d, mDefersToCreateVideoTrack: %d, mDefersToCreateAudioTrack: %d",
140 mProbePkts, mEOF, mFormatCtx->pb ? mFormatCtx->pb->error : 0, mDefersToCreateVideoTrack, mDefersToCreateAudioTrack);
145 FFmpegExtractor::~FFmpegExtractor() {
146 ALOGV("FFmpegExtractor::~FFmpegExtractor");
147 // stop reader here if no track!
150 Mutex::Autolock autoLock(mLock);
154 size_t FFmpegExtractor::countTracks() {
155 return mInitCheck == OK ? mTracks.size() : 0;
158 sp<MediaSource> FFmpegExtractor::getTrack(size_t index) {
159 ALOGV("FFmpegExtractor::getTrack[%d]", index);
161 if (mInitCheck != OK) {
165 if (index >= mTracks.size()) {
169 return new FFmpegSource(this, index);
172 sp<MetaData> FFmpegExtractor::getTrackMetaData(size_t index, uint32_t flags __unused) {
173 ALOGV("FFmpegExtractor::getTrackMetaData[%d]", index);
175 if (mInitCheck != OK) {
179 if (index >= mTracks.size()) {
183 /* Quick and dirty, just get a frame 1/4 in */
184 if (mTracks.itemAt(index).mIndex == mVideoStreamIdx) {
185 int64_t thumb_ts = av_rescale_q((mFormatCtx->streams[mVideoStreamIdx]->duration / 4),
186 mFormatCtx->streams[mVideoStreamIdx]->time_base, AV_TIME_BASE_Q);
187 mTracks.itemAt(index).mMeta->setInt64(kKeyThumbnailTime, thumb_ts);
190 return mTracks.itemAt(index).mMeta;
193 sp<MetaData> FFmpegExtractor::getMetaData() {
194 ALOGV("FFmpegExtractor::getMetaData");
196 if (mInitCheck != OK) {
203 uint32_t FFmpegExtractor::flags() const {
204 ALOGV("FFmpegExtractor::flags");
206 if (mInitCheck != OK) {
210 uint32_t flags = CAN_PAUSE;
212 if (mFormatCtx->duration != AV_NOPTS_VALUE) {
213 flags |= CAN_SEEK_BACKWARD | CAN_SEEK_FORWARD | CAN_SEEK;
219 int FFmpegExtractor::check_extradata(AVCodecContext *avctx)
221 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
222 const char *name = NULL;
223 bool *defersToCreateTrack = NULL;
224 AVBitStreamFilterContext **bsfc = NULL;
227 if (avctx->codec_type == AVMEDIA_TYPE_VIDEO) {
229 defersToCreateTrack = &mDefersToCreateVideoTrack;
230 } else if (avctx->codec_type == AVMEDIA_TYPE_AUDIO){
232 defersToCreateTrack = &mDefersToCreateAudioTrack;
235 codec_id = avctx->codec_id;
238 if (codec_id != AV_CODEC_ID_H264
239 && codec_id != AV_CODEC_ID_MPEG4
240 && codec_id != AV_CODEC_ID_MPEG1VIDEO
241 && codec_id != AV_CODEC_ID_MPEG2VIDEO
242 && codec_id != AV_CODEC_ID_AAC) {
246 // is extradata compatible with android?
247 if (codec_id != AV_CODEC_ID_AAC) {
248 int is_compatible = is_extradata_compatible_with_android(avctx);
249 if (!is_compatible) {
250 ALOGI("%s extradata is not compatible with android, should to extract it from bitstream",
251 av_get_media_type_string(avctx->codec_type));
252 *defersToCreateTrack = true;
253 *bsfc = NULL; // H264 don't need bsfc, only AAC?
259 if (codec_id == AV_CODEC_ID_AAC) {
260 name = "aac_adtstoasc";
263 if (avctx->extradata_size <= 0) {
264 ALOGI("No %s extradata found, should to extract it from bitstream",
265 av_get_media_type_string(avctx->codec_type));
266 *defersToCreateTrack = true;
267 //CHECK(name != NULL);
268 if (!*bsfc && name) {
269 *bsfc = av_bitstream_filter_init(name);
271 ALOGE("Cannot open the %s BSF!", name);
272 *defersToCreateTrack = false;
275 ALOGV("open the %s bsf", name);
285 void FFmpegExtractor::printTime(int64_t time)
287 int hours, mins, secs, us;
289 if (time == AV_NOPTS_VALUE)
292 secs = time / AV_TIME_BASE;
293 us = time % AV_TIME_BASE;
298 ALOGI("the time is %02d:%02d:%02d.%02d",
299 hours, mins, secs, (100 * us) / AV_TIME_BASE);
302 bool FFmpegExtractor::is_codec_supported(enum AVCodecID codec_id)
304 bool supported = false;
307 case AV_CODEC_ID_H264:
308 case AV_CODEC_ID_MPEG4:
309 case AV_CODEC_ID_H263:
310 case AV_CODEC_ID_H263P:
311 case AV_CODEC_ID_H263I:
312 case AV_CODEC_ID_AAC:
313 case AV_CODEC_ID_AC3:
314 case AV_CODEC_ID_MP2:
315 case AV_CODEC_ID_MP3:
316 case AV_CODEC_ID_MPEG1VIDEO:
317 case AV_CODEC_ID_MPEG2VIDEO:
318 case AV_CODEC_ID_WMV1:
319 case AV_CODEC_ID_WMV2:
320 case AV_CODEC_ID_WMV3:
321 case AV_CODEC_ID_VC1:
322 case AV_CODEC_ID_WMAV1:
323 case AV_CODEC_ID_WMAV2:
324 case AV_CODEC_ID_WMAPRO:
325 case AV_CODEC_ID_WMALOSSLESS:
326 case AV_CODEC_ID_RV20:
327 case AV_CODEC_ID_RV30:
328 case AV_CODEC_ID_RV40:
329 case AV_CODEC_ID_COOK:
330 case AV_CODEC_ID_APE:
331 case AV_CODEC_ID_DTS:
332 case AV_CODEC_ID_FLAC:
333 case AV_CODEC_ID_FLV1:
334 case AV_CODEC_ID_VORBIS:
335 case AV_CODEC_ID_HEVC:
340 ALOGD("unsuppoted codec(%s), but give it a chance",
341 avcodec_get_name(codec_id));
342 //Won't promise that the following codec id can be supported.
343 //Just give these codecs a chance.
351 sp<MetaData> FFmpegExtractor::setVideoFormat(AVStream *stream)
353 AVCodecContext *avctx = NULL;
354 sp<MetaData> meta = NULL;
356 avctx = stream->codec;
357 CHECK_EQ(avctx->codec_type, AVMEDIA_TYPE_VIDEO);
359 switch(avctx->codec_id) {
360 case AV_CODEC_ID_H264:
361 if (avctx->extradata[0] == 1) {
362 meta = setAVCFormat(avctx);
364 meta = setH264Format(avctx);
367 case AV_CODEC_ID_MPEG4:
368 meta = setMPEG4Format(avctx);
370 case AV_CODEC_ID_H263:
371 case AV_CODEC_ID_H263P:
372 case AV_CODEC_ID_H263I:
373 meta = setH263Format(avctx);
375 case AV_CODEC_ID_MPEG1VIDEO:
376 case AV_CODEC_ID_MPEG2VIDEO:
377 meta = setMPEG2VIDEOFormat(avctx);
379 case AV_CODEC_ID_VC1:
380 meta = setVC1Format(avctx);
382 case AV_CODEC_ID_WMV1:
383 meta = setWMV1Format(avctx);
385 case AV_CODEC_ID_WMV2:
386 meta = setWMV2Format(avctx);
388 case AV_CODEC_ID_WMV3:
389 meta = setWMV3Format(avctx);
391 case AV_CODEC_ID_RV20:
392 meta = setRV20Format(avctx);
394 case AV_CODEC_ID_RV30:
395 meta = setRV30Format(avctx);
397 case AV_CODEC_ID_RV40:
398 meta = setRV40Format(avctx);
400 case AV_CODEC_ID_FLV1:
401 meta = setFLV1Format(avctx);
403 case AV_CODEC_ID_HEVC:
404 meta = setHEVCFormat(avctx);
406 case AV_CODEC_ID_VP8:
407 meta = setVP8Format(avctx);
409 case AV_CODEC_ID_VP9:
410 meta = setVP9Format(avctx);
413 ALOGD("unsuppoted video codec(id:%d, name:%s), but give it a chance",
414 avctx->codec_id, avcodec_get_name(avctx->codec_id));
416 meta->setInt32(kKeyCodecId, avctx->codec_id);
417 meta->setCString(kKeyMIMEType, MEDIA_MIMETYPE_VIDEO_FFMPEG);
418 if (avctx->extradata_size > 0) {
419 meta->setData(kKeyRawCodecSpecificData, 0, avctx->extradata, avctx->extradata_size);
421 //CHECK(!"Should not be here. Unsupported codec.");
426 ALOGI("width: %d, height: %d, bit_rate: %d",
427 avctx->width, avctx->height, avctx->bit_rate);
429 meta->setInt32(kKeyWidth, avctx->width);
430 meta->setInt32(kKeyHeight, avctx->height);
431 if (avctx->bit_rate > 0) {
432 meta->setInt32(kKeyBitRate, avctx->bit_rate);
434 setDurationMetaData(stream, meta);
440 sp<MetaData> FFmpegExtractor::setAudioFormat(AVStream *stream)
442 AVCodecContext *avctx = NULL;
443 sp<MetaData> meta = NULL;
445 avctx = stream->codec;
446 CHECK_EQ(avctx->codec_type, AVMEDIA_TYPE_AUDIO);
448 switch(avctx->codec_id) {
449 case AV_CODEC_ID_MP2:
450 meta = setMP2Format(avctx);
452 case AV_CODEC_ID_MP3:
453 meta = setMP3Format(avctx);
455 case AV_CODEC_ID_VORBIS:
456 meta = setVORBISFormat(avctx);
458 case AV_CODEC_ID_AC3:
459 meta = setAC3Format(avctx);
461 case AV_CODEC_ID_AAC:
462 meta = setAACFormat(avctx);
464 case AV_CODEC_ID_WMAV1:
465 meta = setWMAV1Format(avctx);
467 case AV_CODEC_ID_WMAV2:
468 meta = setWMAV2Format(avctx);
470 case AV_CODEC_ID_WMAPRO:
471 meta = setWMAProFormat(avctx);
473 case AV_CODEC_ID_WMALOSSLESS:
474 meta = setWMALossLessFormat(avctx);
476 case AV_CODEC_ID_COOK:
477 meta = setRAFormat(avctx);
479 case AV_CODEC_ID_APE:
480 meta = setAPEFormat(avctx);
482 case AV_CODEC_ID_DTS:
483 meta = setDTSFormat(avctx);
485 case AV_CODEC_ID_FLAC:
486 meta = setFLACFormat(avctx);
489 ALOGD("unsuppoted audio codec(id:%d, name:%s), but give it a chance",
490 avctx->codec_id, avcodec_get_name(avctx->codec_id));
492 meta->setInt32(kKeyCodecId, avctx->codec_id);
493 meta->setInt32(kKeyCodedSampleBits, avctx->bits_per_coded_sample);
494 meta->setCString(kKeyMIMEType, MEDIA_MIMETYPE_AUDIO_FFMPEG);
495 if (avctx->extradata_size > 0) {
496 meta->setData(kKeyRawCodecSpecificData, 0, avctx->extradata, avctx->extradata_size);
498 //CHECK(!"Should not be here. Unsupported codec.");
503 ALOGD("bit_rate: %d, sample_rate: %d, channels: %d, "
504 "bits_per_coded_sample: %d, block_align: %d "
505 "bits_per_raw_sample: %d, sample_format: %d",
506 avctx->bit_rate, avctx->sample_rate, avctx->channels,
507 avctx->bits_per_coded_sample, avctx->block_align,
508 avctx->bits_per_raw_sample, avctx->sample_fmt);
510 meta->setInt32(kKeyChannelCount, avctx->channels);
511 meta->setInt32(kKeyBitRate, avctx->bit_rate);
512 meta->setInt32(kKeyBitsPerSample, av_get_bytes_per_sample(avctx->sample_fmt) > 2 ? 24 : 16);
513 meta->setInt32(kKeySampleRate, avctx->sample_rate);
514 meta->setInt32(kKeyBlockAlign, avctx->block_align);
515 meta->setInt32(kKeySampleFormat, avctx->sample_fmt);
516 setDurationMetaData(stream, meta);
522 void FFmpegExtractor::setDurationMetaData(AVStream *stream, sp<MetaData> &meta)
524 AVCodecContext *avctx = stream->codec;
526 if (stream->duration != AV_NOPTS_VALUE) {
527 int64_t duration = av_rescale_q(stream->duration, stream->time_base, AV_TIME_BASE_Q);
529 const char *s = av_get_media_type_string(avctx->codec_type);
530 if (stream->start_time != AV_NOPTS_VALUE) {
531 ALOGV("%s startTime:%lld", s, stream->start_time);
533 ALOGV("%s startTime:N/A", s);
535 meta->setInt64(kKeyDuration, duration);
537 // default when no stream duration
538 meta->setInt64(kKeyDuration, mFormatCtx->duration);
542 int FFmpegExtractor::stream_component_open(int stream_index)
544 TrackInfo *trackInfo = NULL;
545 AVCodecContext *avctx = NULL;
546 sp<MetaData> meta = NULL;
547 bool supported = false;
549 const void *data = NULL;
553 ALOGI("stream_index: %d", stream_index);
554 if (stream_index < 0 || stream_index >= (int)mFormatCtx->nb_streams)
556 avctx = mFormatCtx->streams[stream_index]->codec;
558 supported = is_codec_supported(avctx->codec_id);
561 ALOGE("unsupport the codec(%s)", avcodec_get_name(avctx->codec_id));
564 ALOGI("support the codec(%s)", avcodec_get_name(avctx->codec_id));
567 for (size_t i = 0; i < mTracks.size(); ++i) {
568 if (stream_index == mTracks.editItemAt(i).mIndex) {
569 ALOGE("this track already exists");
574 mFormatCtx->streams[stream_index]->discard = AVDISCARD_DEFAULT;
577 av_get_codec_tag_string(tagbuf, sizeof(tagbuf), avctx->codec_tag);
578 ALOGV("Tag %s/0x%08x with codec(%s)\n", tagbuf, avctx->codec_tag, avcodec_get_name(avctx->codec_id));
580 switch (avctx->codec_type) {
581 case AVMEDIA_TYPE_VIDEO:
582 if (mVideoStreamIdx == -1)
583 mVideoStreamIdx = stream_index;
584 if (mVideoStream == NULL)
585 mVideoStream = mFormatCtx->streams[stream_index];
587 ret = check_extradata(avctx);
590 // disable the stream
591 mVideoStreamIdx = -1;
593 packet_queue_flush(&mVideoQ);
594 mFormatCtx->streams[stream_index]->discard = AVDISCARD_ALL;
599 if (avctx->extradata) {
600 ALOGV("video stream extradata:");
601 hexdump(avctx->extradata, avctx->extradata_size);
603 ALOGV("video stream no extradata, but we can ignore it.");
606 meta = setVideoFormat(mVideoStream);
608 ALOGE("setVideoFormat failed");
612 ALOGV("create a video track");
614 trackInfo = &mTracks.editItemAt(mTracks.size() - 1);
615 trackInfo->mIndex = stream_index;
616 trackInfo->mMeta = meta;
617 trackInfo->mStream = mVideoStream;
618 trackInfo->mQueue = &mVideoQ;
620 mDefersToCreateVideoTrack = false;
623 case AVMEDIA_TYPE_AUDIO:
624 if (mAudioStreamIdx == -1)
625 mAudioStreamIdx = stream_index;
626 if (mAudioStream == NULL)
627 mAudioStream = mFormatCtx->streams[stream_index];
629 ret = check_extradata(avctx);
632 // disable the stream
633 mAudioStreamIdx = -1;
635 packet_queue_flush(&mAudioQ);
636 mFormatCtx->streams[stream_index]->discard = AVDISCARD_ALL;
641 if (avctx->extradata) {
642 ALOGV("audio stream extradata(%d):", avctx->extradata_size);
643 hexdump(avctx->extradata, avctx->extradata_size);
645 ALOGV("audio stream no extradata, but we can ignore it.");
648 meta = setAudioFormat(mAudioStream);
650 ALOGE("setAudioFormat failed");
654 ALOGV("create a audio track");
656 trackInfo = &mTracks.editItemAt(mTracks.size() - 1);
657 trackInfo->mIndex = stream_index;
658 trackInfo->mMeta = meta;
659 trackInfo->mStream = mAudioStream;
660 trackInfo->mQueue = &mAudioQ;
662 mDefersToCreateAudioTrack = false;
665 case AVMEDIA_TYPE_SUBTITLE:
667 CHECK(!"Should not be here. Unsupported media type.");
670 CHECK(!"Should not be here. Unsupported media type.");
676 void FFmpegExtractor::stream_component_close(int stream_index)
678 AVCodecContext *avctx;
680 if (stream_index < 0 || stream_index >= (int)mFormatCtx->nb_streams)
682 avctx = mFormatCtx->streams[stream_index]->codec;
684 switch (avctx->codec_type) {
685 case AVMEDIA_TYPE_VIDEO:
686 ALOGV("packet_queue_abort videoq");
687 packet_queue_abort(&mVideoQ);
688 ALOGV("packet_queue_end videoq");
689 packet_queue_flush(&mVideoQ);
691 case AVMEDIA_TYPE_AUDIO:
692 ALOGV("packet_queue_abort audioq");
693 packet_queue_abort(&mAudioQ);
694 ALOGV("packet_queue_end audioq");
695 packet_queue_flush(&mAudioQ);
697 case AVMEDIA_TYPE_SUBTITLE:
703 mFormatCtx->streams[stream_index]->discard = AVDISCARD_ALL;
704 switch (avctx->codec_type) {
705 case AVMEDIA_TYPE_VIDEO:
707 mVideoStreamIdx = -1;
709 av_bitstream_filter_close(mVideoBsfc);
713 case AVMEDIA_TYPE_AUDIO:
715 mAudioStreamIdx = -1;
717 av_bitstream_filter_close(mAudioBsfc);
721 case AVMEDIA_TYPE_SUBTITLE:
728 void FFmpegExtractor::reachedEOS(enum AVMediaType media_type)
730 Mutex::Autolock autoLock(mLock);
732 if (media_type == AVMEDIA_TYPE_VIDEO) {
733 mVideoEOSReceived = true;
734 } else if (media_type == AVMEDIA_TYPE_AUDIO) {
735 mAudioEOSReceived = true;
740 /* seek in the stream */
741 int FFmpegExtractor::stream_seek(int64_t pos, enum AVMediaType media_type,
742 MediaSource::ReadOptions::SeekMode mode)
744 Mutex::Autolock _l(mLock);
746 if (mSeekIdx >= 0 || (mVideoStreamIdx >= 0
747 && mAudioStreamIdx >= 0
748 && media_type == AVMEDIA_TYPE_AUDIO
749 && !mVideoEOSReceived)) {
754 if (mAudioStreamIdx >= 0)
755 packet_queue_flush(&mAudioQ);
756 if (mVideoStreamIdx >= 0)
757 packet_queue_flush(&mVideoQ);
759 mSeekIdx = media_type == AVMEDIA_TYPE_VIDEO ? mVideoStreamIdx : mAudioStreamIdx;
762 //mSeekFlags &= ~AVSEEK_FLAG_BYTE;
763 //if (mSeekByBytes) {
764 // mSeekFlags |= AVSEEK_FLAG_BYTE;
768 case MediaSource::ReadOptions::SEEK_PREVIOUS_SYNC:
769 mSeekMin = INT64_MIN;
772 case MediaSource::ReadOptions::SEEK_NEXT_SYNC:
774 mSeekMax = INT64_MAX;
776 case MediaSource::ReadOptions::SEEK_CLOSEST_SYNC:
777 mSeekMin = INT64_MIN;
778 mSeekMax = INT64_MAX;
780 case MediaSource::ReadOptions::SEEK_CLOSEST:
781 mSeekMin = INT64_MIN;
788 mCondition.wait(mLock);
793 int FFmpegExtractor::decode_interrupt_cb(void *ctx)
795 FFmpegExtractor *extractor = static_cast<FFmpegExtractor *>(ctx);
796 return extractor->mAbortRequest;
799 void FFmpegExtractor::fetchStuffsFromSniffedMeta(const sp<AMessage> &meta)
805 CHECK(meta->findString("extended-extractor-url", &url));
806 CHECK(url.c_str() != NULL);
807 CHECK(url.size() < PATH_MAX);
809 memcpy(mFilename, url.c_str(), url.size());
810 mFilename[url.size()] = '\0';
813 CHECK(meta->findString("extended-extractor-mime", &mime));
814 CHECK(mime.c_str() != NULL);
815 mMeta->setCString(kKeyMIMEType, mime.c_str());
818 void FFmpegExtractor::setFFmpegDefaultOpts()
821 #if DEBUG_DISABLE_VIDEO
826 #if DEBUG_DISABLE_AUDIO
832 mSeekByBytes = 0; /* seek by bytes 0=off 1=on -1=auto" */
833 mDuration = AV_NOPTS_VALUE;
834 mSeekPos = AV_NOPTS_VALUE;
835 mSeekMin = INT64_MIN;
836 mSeekMax = INT64_MAX;
839 mVideoStreamIdx = -1;
840 mAudioStreamIdx = -1;
843 mDefersToCreateVideoTrack = false;
844 mDefersToCreateAudioTrack = false;
857 int FFmpegExtractor::initStreams()
861 status_t status = UNKNOWN_ERROR;
863 int ret = 0, audio_ret = -1, video_ret = -1;
864 int pkt_in_play_range = 0;
865 AVDictionaryEntry *t = NULL;
866 AVDictionary **opts = NULL;
867 int orig_nb_streams = 0;
868 int st_index[AVMEDIA_TYPE_NB] = {0};
869 int wanted_stream[AVMEDIA_TYPE_NB] = {0};
870 st_index[AVMEDIA_TYPE_AUDIO] = -1;
871 st_index[AVMEDIA_TYPE_VIDEO] = -1;
872 wanted_stream[AVMEDIA_TYPE_AUDIO] = -1;
873 wanted_stream[AVMEDIA_TYPE_VIDEO] = -1;
874 const char *mime = NULL;
876 setFFmpegDefaultOpts();
878 status = initFFmpeg();
883 mFFmpegInited = true;
885 mFormatCtx = avformat_alloc_context();
888 ALOGE("oom for alloc avformat context");
892 mFormatCtx->interrupt_callback.callback = decode_interrupt_cb;
893 mFormatCtx->interrupt_callback.opaque = this;
894 ALOGV("mFilename: %s", mFilename);
895 err = avformat_open_input(&mFormatCtx, mFilename, NULL, &format_opts);
897 ALOGE("%s: avformat_open_input failed, err:%s", mFilename, av_err2str(err));
902 if ((t = av_dict_get(format_opts, "", NULL, AV_DICT_IGNORE_SUFFIX))) {
903 ALOGE("Option %s not found.\n", t->key);
904 //ret = AVERROR_OPTION_NOT_FOUND;
910 mFormatCtx->flags |= AVFMT_FLAG_GENPTS;
912 opts = setup_find_stream_info_opts(mFormatCtx, codec_opts);
913 orig_nb_streams = mFormatCtx->nb_streams;
915 err = avformat_find_stream_info(mFormatCtx, opts);
917 ALOGE("%s: could not find stream info, err:%s", mFilename, av_err2str(err));
921 for (i = 0; i < orig_nb_streams; i++)
922 av_dict_free(&opts[i]);
926 mFormatCtx->pb->eof_reached = 0; // FIXME hack, ffplay maybe should not use url_feof() to test for the end
928 if (mSeekByBytes < 0)
929 mSeekByBytes = !!(mFormatCtx->iformat->flags & AVFMT_TS_DISCONT)
930 && strcmp("ogg", mFormatCtx->iformat->name);
932 for (i = 0; i < (int)mFormatCtx->nb_streams; i++)
933 mFormatCtx->streams[i]->discard = AVDISCARD_ALL;
935 st_index[AVMEDIA_TYPE_VIDEO] =
936 av_find_best_stream(mFormatCtx, AVMEDIA_TYPE_VIDEO,
937 wanted_stream[AVMEDIA_TYPE_VIDEO], -1, NULL, 0);
939 st_index[AVMEDIA_TYPE_AUDIO] =
940 av_find_best_stream(mFormatCtx, AVMEDIA_TYPE_AUDIO,
941 wanted_stream[AVMEDIA_TYPE_AUDIO],
942 st_index[AVMEDIA_TYPE_VIDEO],
945 av_dump_format(mFormatCtx, 0, mFilename, 0);
948 if (mFormatCtx->duration != AV_NOPTS_VALUE &&
949 mFormatCtx->start_time != AV_NOPTS_VALUE) {
950 int hours, mins, secs, us;
952 ALOGV("file startTime: %lld", mFormatCtx->start_time);
954 mDuration = mFormatCtx->duration;
956 secs = mDuration / AV_TIME_BASE;
957 us = mDuration % AV_TIME_BASE;
962 ALOGI("the duration is %02d:%02d:%02d.%02d",
963 hours, mins, secs, (100 * us) / AV_TIME_BASE);
966 packet_queue_init(&mVideoQ);
967 packet_queue_init(&mAudioQ);
969 if (st_index[AVMEDIA_TYPE_AUDIO] >= 0) {
970 audio_ret = stream_component_open(st_index[AVMEDIA_TYPE_AUDIO]);
972 packet_queue_start(&mAudioQ);
975 if (st_index[AVMEDIA_TYPE_VIDEO] >= 0) {
976 video_ret = stream_component_open(st_index[AVMEDIA_TYPE_VIDEO]);
978 packet_queue_start(&mVideoQ);
981 if ( audio_ret < 0 && video_ret < 0) {
982 ALOGE("%s: could not open codecs\n", mFilename);
993 void FFmpegExtractor::deInitStreams()
995 packet_queue_destroy(&mVideoQ);
996 packet_queue_destroy(&mAudioQ);
999 avformat_close_input(&mFormatCtx);
1002 if (mFFmpegInited) {
1007 status_t FFmpegExtractor::startReaderThread() {
1008 ALOGV("Starting reader thread");
1010 if (mReaderThreadStarted)
1013 pthread_attr_t attr;
1014 pthread_attr_init(&attr);
1015 pthread_attr_setdetachstate(&attr, PTHREAD_CREATE_JOINABLE);
1017 ALOGD("Reader thread starting");
1019 pthread_create(&mReaderThread, &attr, ReaderWrapper, this);
1020 pthread_attr_destroy(&attr);
1022 mReaderThreadStarted = true;
1023 mCondition.signal();
1028 void FFmpegExtractor::stopReaderThread() {
1029 ALOGV("Stopping reader thread");
1033 if (!mReaderThreadStarted) {
1034 ALOGD("Reader thread have been stopped");
1040 mCondition.signal();
1042 /* close each stream */
1043 if (mAudioStreamIdx >= 0)
1044 stream_component_close(mAudioStreamIdx);
1045 if (mVideoStreamIdx >= 0)
1046 stream_component_close(mVideoStreamIdx);
1049 pthread_join(mReaderThread, NULL);
1053 avformat_close_input(&mFormatCtx);
1056 mReaderThreadStarted = false;
1057 ALOGD("Reader thread stopped");
1063 void *FFmpegExtractor::ReaderWrapper(void *me) {
1064 ((FFmpegExtractor *)me)->readerEntry();
1069 void FFmpegExtractor::readerEntry() {
1071 AVPacket pkt1, *pkt = &pkt1;
1073 int pkt_in_play_range = 0;
1077 pid_t tid = gettid();
1078 androidSetThreadPriority(tid,
1079 mVideoStreamIdx >= 0 ? ANDROID_PRIORITY_NORMAL : ANDROID_PRIORITY_AUDIO);
1080 prctl(PR_SET_NAME, (unsigned long)"FFmpegExtractor Thread", 0, 0, 0);
1082 ALOGV("FFmpegExtractor wait for signal");
1083 while (!mReaderThreadStarted && !mAbortRequest) {
1084 mCondition.wait(mLock);
1086 ALOGV("FFmpegExtractor ready to run");
1088 if (mAbortRequest) {
1092 mVideoEOSReceived = false;
1093 mAudioEOSReceived = false;
1095 while (!mAbortRequest) {
1097 if (mPaused != mLastPaused) {
1098 mLastPaused = mPaused;
1100 mReadPauseReturn = av_read_pause(mFormatCtx);
1102 av_read_play(mFormatCtx);
1104 #if CONFIG_RTSP_DEMUXER || CONFIG_MMSH_PROTOCOL
1106 (!strcmp(mFormatCtx->iformat->name, "rtsp") ||
1107 (mFormatCtx->pb && !strncmp(mFilename, "mmsh:", 5)))) {
1108 /* wait 10 ms to avoid trying to get another packet */
1115 if (mSeekIdx >= 0) {
1116 Mutex::Autolock _l(mLock);
1117 ALOGV("readerEntry, mSeekIdx: %d mSeekPos: %lld (%lld/%lld)", mSeekIdx, mSeekPos, mSeekMin, mSeekMax);
1118 ret = avformat_seek_file(mFormatCtx, -1, mSeekMin, mSeekPos, mSeekMax, 0);
1120 ALOGE("%s: error while seeking", mFormatCtx->filename);
1122 if (mAudioStreamIdx >= 0) {
1123 packet_queue_flush(&mAudioQ);
1124 packet_queue_put(&mAudioQ, &mAudioQ.flush_pkt);
1126 if (mVideoStreamIdx >= 0) {
1127 packet_queue_flush(&mVideoQ);
1128 packet_queue_put(&mVideoQ, &mVideoQ.flush_pkt);
1133 mCondition.signal();
1136 /* if the queue are full, no need to read more */
1137 if ( mAudioQ.size + mVideoQ.size > MAX_QUEUE_SIZE
1138 || ( (mAudioQ .size > MIN_AUDIOQ_SIZE || mAudioStreamIdx < 0)
1139 && (mVideoQ .nb_packets > MIN_FRAMES || mVideoStreamIdx < 0))) {
1140 #if DEBUG_READ_ENTRY
1141 ALOGV("readerEntry, full(wtf!!!), mVideoQ.size: %d, mVideoQ.nb_packets: %d, mAudioQ.size: %d, mAudioQ.nb_packets: %d",
1142 mVideoQ.size, mVideoQ.nb_packets, mAudioQ.size, mAudioQ.nb_packets);
1145 mExtractorMutex.lock();
1146 mCondition.waitRelative(mExtractorMutex, milliseconds(10));
1147 mExtractorMutex.unlock();
1152 if (mVideoStreamIdx >= 0) {
1153 packet_queue_put_nullpacket(&mVideoQ, mVideoStreamIdx);
1155 if (mAudioStreamIdx >= 0) {
1156 packet_queue_put_nullpacket(&mAudioQ, mAudioStreamIdx);
1159 mExtractorMutex.lock();
1160 mCondition.waitRelative(mExtractorMutex, milliseconds(10));
1162 mExtractorMutex.unlock();
1166 ret = av_read_frame(mFormatCtx, pkt);
1172 if (mFormatCtx->pb && mFormatCtx->pb->error) {
1173 ALOGE("mFormatCtx->pb->error: %d", mFormatCtx->pb->error);
1177 mExtractorMutex.lock();
1178 mCondition.waitRelative(mExtractorMutex, milliseconds(10));
1179 mExtractorMutex.unlock();
1183 if (pkt->stream_index == mVideoStreamIdx) {
1184 if (mDefersToCreateVideoTrack) {
1185 AVCodecContext *avctx = mFormatCtx->streams[mVideoStreamIdx]->codec;
1187 int i = parser_split(avctx, pkt->data, pkt->size);
1188 if (i > 0 && i < FF_MAX_EXTRADATA_SIZE) {
1189 if (avctx->extradata)
1190 av_freep(&avctx->extradata);
1191 avctx->extradata_size= i;
1192 avctx->extradata = (uint8_t *)av_malloc(avctx->extradata_size + FF_INPUT_BUFFER_PADDING_SIZE);
1193 if (!avctx->extradata) {
1194 //return AVERROR(ENOMEM);
1195 ret = AVERROR(ENOMEM);
1198 // sps + pps(there may be sei in it)
1199 memcpy(avctx->extradata, pkt->data, avctx->extradata_size);
1200 memset(avctx->extradata + i, 0, FF_INPUT_BUFFER_PADDING_SIZE);
1202 av_free_packet(pkt);
1206 stream_component_open(mVideoStreamIdx);
1207 if (!mDefersToCreateVideoTrack)
1208 ALOGI("probe packet counter: %d when create video track ok", mProbePkts);
1209 if (mProbePkts == EXTRACTOR_MAX_PROBE_PACKETS)
1210 ALOGI("probe packet counter to max: %d, create video track: %d",
1211 mProbePkts, !mDefersToCreateVideoTrack);
1213 } else if (pkt->stream_index == mAudioStreamIdx) {
1217 AVCodecContext *avctx = mFormatCtx->streams[mAudioStreamIdx]->codec;
1218 if (mAudioBsfc && pkt && pkt->data) {
1219 ret = av_bitstream_filter_filter(mAudioBsfc, avctx, NULL, &outbuf, &outbuf_size,
1220 pkt->data, pkt->size, pkt->flags & AV_PKT_FLAG_KEY);
1222 if (ret < 0 ||!outbuf_size) {
1223 av_free_packet(pkt);
1226 if (outbuf && outbuf != pkt->data) {
1227 memmove(pkt->data, outbuf, outbuf_size);
1228 pkt->size = outbuf_size;
1231 if (mDefersToCreateAudioTrack) {
1232 if (avctx->extradata_size <= 0) {
1233 av_free_packet(pkt);
1236 stream_component_open(mAudioStreamIdx);
1237 if (!mDefersToCreateAudioTrack)
1238 ALOGI("probe packet counter: %d when create audio track ok", mProbePkts);
1239 if (mProbePkts == EXTRACTOR_MAX_PROBE_PACKETS)
1240 ALOGI("probe packet counter to max: %d, create audio track: %d",
1241 mProbePkts, !mDefersToCreateAudioTrack);
1245 if (pkt->stream_index == mAudioStreamIdx) {
1246 packet_queue_put(&mAudioQ, pkt);
1247 } else if (pkt->stream_index == mVideoStreamIdx) {
1248 packet_queue_put(&mVideoQ, pkt);
1250 av_free_packet(pkt);
1257 ALOGV("FFmpegExtractor exit thread(readerEntry)");
1260 ////////////////////////////////////////////////////////////////////////////////
1262 FFmpegSource::FFmpegSource(
1263 const sp<FFmpegExtractor> &extractor, size_t index)
1264 : mExtractor(extractor),
1268 mStream(mExtractor->mTracks.itemAt(index).mStream),
1269 mQueue(mExtractor->mTracks.itemAt(index).mQueue),
1270 mLastPTS(AV_NOPTS_VALUE),
1271 mTargetTime(AV_NOPTS_VALUE) {
1272 sp<MetaData> meta = mExtractor->mTracks.itemAt(index).mMeta;
1275 AVCodecContext *avctx = mStream->codec;
1277 /* Parse codec specific data */
1278 if (avctx->codec_id == AV_CODEC_ID_H264
1279 && avctx->extradata_size > 0
1280 && avctx->extradata[0] == 1) {
1286 CHECK(meta->findData(kKeyAVCC, &type, &data, &size));
1288 const uint8_t *ptr = (const uint8_t *)data;
1291 CHECK_EQ((unsigned)ptr[0], 1u); // configurationVersion == 1
1293 // The number of bytes used to encode the length of a NAL unit.
1294 mNALLengthSize = 1 + (ptr[4] & 3);
1296 ALOGV("the stream is AVC, the length of a NAL unit: %d", mNALLengthSize);
1302 mMediaType = mStream->codec->codec_type;
1303 mFirstKeyPktTimestamp = AV_NOPTS_VALUE;
1306 FFmpegSource::~FFmpegSource() {
1307 ALOGV("FFmpegSource::~FFmpegSource %s",
1308 av_get_media_type_string(mMediaType));
1312 status_t FFmpegSource::start(MetaData *params __unused) {
1313 ALOGV("FFmpegSource::start %s",
1314 av_get_media_type_string(mMediaType));
1318 status_t FFmpegSource::stop() {
1319 ALOGV("FFmpegSource::stop %s",
1320 av_get_media_type_string(mMediaType));
1324 sp<MetaData> FFmpegSource::getFormat() {
1325 return mExtractor->mTracks.itemAt(mTrackIndex).mMeta;;
1328 status_t FFmpegSource::read(
1329 MediaBuffer **buffer, const ReadOptions *options) {
1333 bool seeking = false;
1334 bool waitKeyPkt = false;
1335 ReadOptions::SeekMode mode;
1336 int64_t pktTS = AV_NOPTS_VALUE;
1337 int64_t seekTimeUs = AV_NOPTS_VALUE;
1338 int64_t timeUs = AV_NOPTS_VALUE;
1340 status_t status = OK;
1342 int64_t startTimeUs = mStream->start_time == AV_NOPTS_VALUE ? 0 :
1343 av_rescale_q(mStream->start_time, mStream->time_base, AV_TIME_BASE_Q);
1345 if (options && options->getSeekTo(&seekTimeUs, &mode)) {
1346 int64_t seekPTS = seekTimeUs;
1347 ALOGV("~~~%s seekTimeUs: %lld, seekPTS: %lld, mode: %d", av_get_media_type_string(mMediaType), seekTimeUs, seekPTS, mode);
1348 /* add the stream start time */
1349 if (mStream->start_time != AV_NOPTS_VALUE) {
1350 seekPTS += startTimeUs;
1352 ALOGV("~~~%s seekTimeUs[+startTime]: %lld, mode: %d start_time=%lld", av_get_media_type_string(mMediaType), seekPTS, mode, startTimeUs);
1353 seeking = (mExtractor->stream_seek(seekPTS, mMediaType, mode) == SEEK);
1357 if (packet_queue_get(mQueue, &pkt, 1) < 0) {
1358 ALOGD("read %s abort reqeust", av_get_media_type_string(mMediaType));
1359 mExtractor->reachedEOS(mMediaType);
1360 return ERROR_END_OF_STREAM;
1364 if (pkt.data != mQueue->flush_pkt.data) {
1365 av_free_packet(&pkt);
1369 #if WAIT_KEY_PACKET_AFTER_SEEK
1375 if (pkt.data == mQueue->flush_pkt.data) {
1376 ALOGV("read %s flush pkt", av_get_media_type_string(mMediaType));
1377 av_free_packet(&pkt);
1378 mFirstKeyPktTimestamp = AV_NOPTS_VALUE;
1380 } else if (pkt.data == NULL && pkt.size == 0) {
1381 ALOGD("read %s eos pkt", av_get_media_type_string(mMediaType));
1382 av_free_packet(&pkt);
1383 mExtractor->reachedEOS(mMediaType);
1384 return ERROR_END_OF_STREAM;
1387 key = pkt.flags & AV_PKT_FLAG_KEY ? 1 : 0;
1388 pktTS = pkt.pts == AV_NOPTS_VALUE ? pkt.dts : pkt.pts;
1392 ALOGV("drop the non-key packet");
1393 av_free_packet(&pkt);
1396 ALOGV("~~~~~~ got the key packet");
1401 if (pktTS != AV_NOPTS_VALUE && mFirstKeyPktTimestamp == AV_NOPTS_VALUE) {
1402 // update the first key timestamp
1403 mFirstKeyPktTimestamp = pktTS;
1406 MediaBuffer *mediaBuffer = new MediaBuffer(pkt.size + FF_INPUT_BUFFER_PADDING_SIZE);
1407 mediaBuffer->meta_data()->clear();
1408 mediaBuffer->set_range(0, pkt.size);
1411 if (mIsAVC && mNal2AnnexB) {
1412 /* This only works for NAL sizes 3-4 */
1413 CHECK(mNALLengthSize == 3 || mNALLengthSize == 4);
1415 uint8_t *dst = (uint8_t *)mediaBuffer->data();
1416 /* Convert H.264 NAL format to annex b */
1417 status = convertNal2AnnexB(dst, pkt.size, pkt.data, pkt.size, mNALLengthSize);
1419 ALOGE("convertNal2AnnexB failed");
1420 mediaBuffer->release();
1422 av_free_packet(&pkt);
1423 return ERROR_MALFORMED;
1426 memcpy(mediaBuffer->data(), pkt.data, pkt.size);
1429 if (pktTS != AV_NOPTS_VALUE)
1430 timeUs = av_rescale_q(pktTS, mStream->time_base, AV_TIME_BASE_Q) - startTimeUs;
1432 timeUs = SF_NOPTS_VALUE; //FIXME AV_NOPTS_VALUE is negative, but stagefright need positive
1434 // predict the next PTS to use for exact-frame seek below
1435 int64_t nextPTS = AV_NOPTS_VALUE;
1436 if (mLastPTS != AV_NOPTS_VALUE && timeUs > mLastPTS) {
1437 nextPTS = timeUs + (timeUs - mLastPTS);
1439 } else if (mLastPTS == AV_NOPTS_VALUE) {
1444 if (pktTS != AV_NOPTS_VALUE)
1445 ALOGV("read %s pkt, size:%d, key:%d, pktPTS: %lld, pts:%lld, dts:%lld, timeUs[-startTime]:%lld us (%.2f secs) start_time=%lld",
1446 av_get_media_type_string(mMediaType), pkt.size, key, pktTS, pkt.pts, pkt.dts, timeUs, timeUs/1E6, startTimeUs);
1448 ALOGV("read %s pkt, size:%d, key:%d, pts:N/A, dts:N/A, timeUs[-startTime]:N/A",
1449 av_get_media_type_string(mMediaType), pkt.size, key);
1452 mediaBuffer->meta_data()->setInt64(kKeyTime, timeUs);
1453 mediaBuffer->meta_data()->setInt32(kKeyIsSyncFrame, key);
1455 // deal with seek-to-exact-frame, we might be off a bit and Stagefright will assert on us
1456 if (seekTimeUs != AV_NOPTS_VALUE && timeUs < seekTimeUs &&
1457 mode == MediaSource::ReadOptions::SEEK_CLOSEST) {
1458 mTargetTime = seekTimeUs;
1459 mediaBuffer->meta_data()->setInt64(kKeyTargetTime, seekTimeUs);
1462 if (mTargetTime != AV_NOPTS_VALUE) {
1463 if (timeUs == mTargetTime) {
1464 mTargetTime = AV_NOPTS_VALUE;
1465 } else if (nextPTS != AV_NOPTS_VALUE && nextPTS > mTargetTime) {
1466 ALOGV("adjust target frame time to %lld", timeUs);
1467 mediaBuffer->meta_data()->setInt64(kKeyTime, mTargetTime);
1468 mTargetTime = AV_NOPTS_VALUE;
1472 *buffer = mediaBuffer;
1474 av_free_packet(&pkt);
1479 ////////////////////////////////////////////////////////////////////////////////
1483 const char *container;
1486 static formatmap FILE_FORMATS[] = {
1487 {"mpeg", MEDIA_MIMETYPE_CONTAINER_MPEG2PS },
1488 {"mpegts", MEDIA_MIMETYPE_CONTAINER_TS },
1489 {"mov,mp4,m4a,3gp,3g2,mj2", MEDIA_MIMETYPE_CONTAINER_MPEG4 },
1490 {"matroska,webm", MEDIA_MIMETYPE_CONTAINER_MATROSKA },
1491 {"asf", MEDIA_MIMETYPE_CONTAINER_ASF },
1492 {"rm", MEDIA_MIMETYPE_CONTAINER_RM },
1493 {"flv", MEDIA_MIMETYPE_CONTAINER_FLV },
1494 {"swf", MEDIA_MIMETYPE_CONTAINER_FLV },
1495 {"avi", MEDIA_MIMETYPE_CONTAINER_AVI },
1496 {"ape", MEDIA_MIMETYPE_CONTAINER_APE },
1497 {"dts", MEDIA_MIMETYPE_CONTAINER_DTS },
1498 {"flac", MEDIA_MIMETYPE_CONTAINER_FLAC },
1499 {"ac3", MEDIA_MIMETYPE_AUDIO_AC3 },
1500 {"mp3", MEDIA_MIMETYPE_AUDIO_MPEG },
1501 {"wav", MEDIA_MIMETYPE_CONTAINER_WAV },
1502 {"ogg", MEDIA_MIMETYPE_CONTAINER_OGG },
1503 {"vc1", MEDIA_MIMETYPE_CONTAINER_VC1 },
1504 {"hevc", MEDIA_MIMETYPE_CONTAINER_HEVC },
1505 {"divx", MEDIA_MIMETYPE_CONTAINER_DIVX },
1508 static AVCodecContext* getCodecContext(AVFormatContext *ic, AVMediaType codec_type)
1510 unsigned int idx = 0;
1511 AVCodecContext *avctx = NULL;
1513 for (idx = 0; idx < ic->nb_streams; idx++) {
1514 if (ic->streams[idx]->disposition & AV_DISPOSITION_ATTACHED_PIC) {
1515 // FFMPEG converts album art to MJPEG, but we don't want to
1516 // include that in the parsing as MJPEG is not supported by
1517 // Android, which forces the media to be extracted by FFMPEG
1518 // while in fact, Android supports it.
1522 avctx = ic->streams[idx]->codec;
1523 if (avctx->codec_type == codec_type) {
1531 static enum AVCodecID getCodecId(AVFormatContext *ic, AVMediaType codec_type)
1533 AVCodecContext *avctx = getCodecContext(ic, codec_type);
1534 return avctx == NULL ? AV_CODEC_ID_NONE : avctx->codec_id;
1537 static bool hasAudioCodecOnly(AVFormatContext *ic)
1539 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
1540 bool haveVideo = false;
1541 bool haveAudio = false;
1543 if (getCodecId(ic, AVMEDIA_TYPE_VIDEO) != AV_CODEC_ID_NONE) {
1546 if (getCodecId(ic, AVMEDIA_TYPE_AUDIO) != AV_CODEC_ID_NONE) {
1550 if (!haveVideo && haveAudio) {
1557 //FIXME all codecs: frameworks/av/media/libstagefright/codecs/*
1558 static bool isCodecSupportedByStagefright(enum AVCodecID codec_id)
1560 bool supported = false;
1564 case AV_CODEC_ID_HEVC:
1565 case AV_CODEC_ID_H264:
1566 case AV_CODEC_ID_MPEG4:
1567 case AV_CODEC_ID_H263:
1568 case AV_CODEC_ID_H263P:
1569 case AV_CODEC_ID_H263I:
1570 case AV_CODEC_ID_VP6:
1571 case AV_CODEC_ID_VP8:
1572 case AV_CODEC_ID_VP9:
1574 case AV_CODEC_ID_AAC:
1575 case AV_CODEC_ID_MP3:
1576 case AV_CODEC_ID_AMR_NB:
1577 case AV_CODEC_ID_AMR_WB:
1578 case AV_CODEC_ID_FLAC:
1579 case AV_CODEC_ID_VORBIS:
1580 case AV_CODEC_ID_PCM_MULAW: //g711
1581 case AV_CODEC_ID_PCM_ALAW: //g711
1582 case AV_CODEC_ID_GSM_MS:
1583 case AV_CODEC_ID_PCM_U8:
1584 case AV_CODEC_ID_PCM_S16LE:
1585 case AV_CODEC_ID_PCM_S24LE:
1593 ALOGD("%ssuppoted codec(%s) by official Stagefright",
1594 (supported ? "" : "un"),
1595 avcodec_get_name(codec_id));
1600 static void adjustMPEG4Confidence(AVFormatContext *ic, float *confidence)
1602 AVDictionary *tags = NULL;
1603 AVDictionaryEntry *tag = NULL;
1604 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
1607 codec_id = getCodecId(ic, AVMEDIA_TYPE_VIDEO);
1608 if (codec_id != AV_CODEC_ID_NONE
1609 && codec_id != AV_CODEC_ID_HEVC
1610 && codec_id != AV_CODEC_ID_H264
1611 && codec_id != AV_CODEC_ID_MPEG4
1612 && codec_id != AV_CODEC_ID_H263
1613 && codec_id != AV_CODEC_ID_H263P
1614 && codec_id != AV_CODEC_ID_H263I) {
1615 //the MEDIA_MIMETYPE_CONTAINER_MPEG4 of confidence is 0.4f
1616 ALOGI("[mp4]video codec(%s), confidence should be larger than MPEG4Extractor",
1617 avcodec_get_name(codec_id));
1618 *confidence = 0.41f;
1621 codec_id = getCodecId(ic, AVMEDIA_TYPE_AUDIO);
1622 if (codec_id != AV_CODEC_ID_NONE
1623 && codec_id != AV_CODEC_ID_MP3
1624 && codec_id != AV_CODEC_ID_AAC
1625 && codec_id != AV_CODEC_ID_AMR_NB
1626 && codec_id != AV_CODEC_ID_AMR_WB) {
1627 ALOGI("[mp4]audio codec(%s), confidence should be larger than MPEG4Extractor",
1628 avcodec_get_name(codec_id));
1629 *confidence = 0.41f;
1633 tags = ic->metadata;
1634 //NOTE: You can use command to show these tags,
1635 //e.g. "ffprobe -show_format 2012.mov"
1636 tag = av_dict_get(tags, "major_brand", NULL, 0);
1641 ALOGV("major_brand tag is:%s", tag->value);
1643 //when MEDIA_MIMETYPE_CONTAINER_MPEG4
1644 //WTF, MPEG4Extractor.cpp can not extractor mov format
1645 //NOTE: isCompatibleBrand(MPEG4Extractor.cpp)
1646 // Won't promise that the following file types can be played.
1647 // Just give these file types a chance.
1648 // FOURCC('q', 't', ' ', ' '), // Apple's QuickTime
1650 if (!strcmp(tag->value, "qt ")) {
1651 ALOGI("[mp4]format is mov, confidence should be larger than mpeg4");
1652 *confidence = 0.41f;
1656 static void adjustMPEG2TSConfidence(AVFormatContext *ic, float *confidence)
1658 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
1660 codec_id = getCodecId(ic, AVMEDIA_TYPE_VIDEO);
1661 if (codec_id != AV_CODEC_ID_NONE
1662 && codec_id != AV_CODEC_ID_H264
1663 && codec_id != AV_CODEC_ID_MPEG4
1664 && codec_id != AV_CODEC_ID_MPEG1VIDEO
1665 && codec_id != AV_CODEC_ID_MPEG2VIDEO) {
1666 //the MEDIA_MIMETYPE_CONTAINER_MPEG2TS of confidence is 0.1f
1667 ALOGI("[mpeg2ts]video codec(%s), confidence should be larger than MPEG2TSExtractor",
1668 avcodec_get_name(codec_id));
1669 *confidence = 0.11f;
1672 codec_id = getCodecId(ic, AVMEDIA_TYPE_AUDIO);
1673 if (codec_id != AV_CODEC_ID_NONE
1674 && codec_id != AV_CODEC_ID_AAC
1675 && codec_id != AV_CODEC_ID_PCM_S16LE
1676 && codec_id != AV_CODEC_ID_PCM_S24LE
1677 && codec_id != AV_CODEC_ID_MP1
1678 && codec_id != AV_CODEC_ID_MP2
1679 && codec_id != AV_CODEC_ID_MP3) {
1680 ALOGI("[mpeg2ts]audio codec(%s), confidence should be larger than MPEG2TSExtractor",
1681 avcodec_get_name(codec_id));
1682 *confidence = 0.11f;
1686 static void adjustMKVConfidence(AVFormatContext *ic, float *confidence)
1688 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
1690 codec_id = getCodecId(ic, AVMEDIA_TYPE_VIDEO);
1691 if (codec_id != AV_CODEC_ID_NONE
1692 && codec_id != AV_CODEC_ID_H264
1693 && codec_id != AV_CODEC_ID_MPEG4
1694 && codec_id != AV_CODEC_ID_VP6
1695 && codec_id != AV_CODEC_ID_VP8
1696 && codec_id != AV_CODEC_ID_VP9) {
1697 //the MEDIA_MIMETYPE_CONTAINER_MATROSKA of confidence is 0.6f
1698 ALOGI("[mkv]video codec(%s), confidence should be larger than MatroskaExtractor",
1699 avcodec_get_name(codec_id));
1700 *confidence = 0.61f;
1703 codec_id = getCodecId(ic, AVMEDIA_TYPE_AUDIO);
1704 if (codec_id != AV_CODEC_ID_NONE
1705 && codec_id != AV_CODEC_ID_AAC
1706 && codec_id != AV_CODEC_ID_MP3
1707 && codec_id != AV_CODEC_ID_VORBIS) {
1708 ALOGI("[mkv]audio codec(%s), confidence should be larger than MatroskaExtractor",
1709 avcodec_get_name(codec_id));
1710 *confidence = 0.61f;
1714 static void adjustCodecConfidence(AVFormatContext *ic, float *confidence)
1716 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
1718 codec_id = getCodecId(ic, AVMEDIA_TYPE_VIDEO);
1719 if (codec_id != AV_CODEC_ID_NONE) {
1720 if (!isCodecSupportedByStagefright(codec_id)) {
1721 *confidence = 0.88f;
1725 codec_id = getCodecId(ic, AVMEDIA_TYPE_AUDIO);
1726 if (codec_id != AV_CODEC_ID_NONE) {
1727 if (!isCodecSupportedByStagefright(codec_id)) {
1728 *confidence = 0.88f;
1732 if (getCodecId(ic, AVMEDIA_TYPE_VIDEO) != AV_CODEC_ID_NONE
1733 && getCodecId(ic, AVMEDIA_TYPE_AUDIO) == AV_CODEC_ID_MP3) {
1734 *confidence = 0.22f; //larger than MP3Extractor
1738 //TODO need more checks
1739 static void adjustConfidenceIfNeeded(const char *mime,
1740 AVFormatContext *ic, float *confidence)
1743 if (!strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MPEG4)) {
1744 adjustMPEG4Confidence(ic, confidence);
1745 } else if (!strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MPEG2TS)) {
1746 adjustMPEG2TSConfidence(ic, confidence);
1747 } else if (!strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MATROSKA)) {
1748 adjustMKVConfidence(ic, confidence);
1753 if (*confidence > 0.08) {
1758 adjustCodecConfidence(ic, confidence);
1761 static void adjustContainerIfNeeded(const char **mime, AVFormatContext *ic)
1763 const char *newMime = *mime;
1764 enum AVCodecID codec_id = AV_CODEC_ID_NONE;
1766 AVCodecContext *avctx = getCodecContext(ic, AVMEDIA_TYPE_VIDEO);
1767 if (avctx != NULL && getDivXVersion(avctx) >= 0) {
1768 newMime = MEDIA_MIMETYPE_VIDEO_DIVX;
1770 } else if (hasAudioCodecOnly(ic)) {
1771 codec_id = getCodecId(ic, AVMEDIA_TYPE_AUDIO);
1772 CHECK(codec_id != AV_CODEC_ID_NONE);
1774 case AV_CODEC_ID_MP3:
1775 newMime = MEDIA_MIMETYPE_AUDIO_MPEG;
1777 case AV_CODEC_ID_AAC:
1778 newMime = MEDIA_MIMETYPE_AUDIO_AAC;
1780 case AV_CODEC_ID_VORBIS:
1781 newMime = MEDIA_MIMETYPE_AUDIO_VORBIS;
1783 case AV_CODEC_ID_FLAC:
1784 newMime = MEDIA_MIMETYPE_AUDIO_FLAC;
1786 case AV_CODEC_ID_AC3:
1787 newMime = MEDIA_MIMETYPE_AUDIO_AC3;
1789 case AV_CODEC_ID_APE:
1790 newMime = MEDIA_MIMETYPE_AUDIO_APE;
1792 case AV_CODEC_ID_DTS:
1793 newMime = MEDIA_MIMETYPE_AUDIO_DTS;
1795 case AV_CODEC_ID_MP2:
1796 newMime = MEDIA_MIMETYPE_AUDIO_MPEG_LAYER_II;
1798 case AV_CODEC_ID_COOK:
1799 newMime = MEDIA_MIMETYPE_AUDIO_RA;
1801 case AV_CODEC_ID_WMAV1:
1802 case AV_CODEC_ID_WMAV2:
1803 case AV_CODEC_ID_WMAPRO:
1804 case AV_CODEC_ID_WMALOSSLESS:
1805 newMime = MEDIA_MIMETYPE_AUDIO_WMA;
1811 if (!strcmp(*mime, MEDIA_MIMETYPE_CONTAINER_FFMPEG)) {
1812 newMime = MEDIA_MIMETYPE_AUDIO_FFMPEG;
1816 if (strcmp(*mime, newMime)) {
1817 ALOGI("adjust mime(%s -> %s)", *mime, newMime);
1822 static const char *findMatchingContainer(const char *name)
1825 #if SUPPOURT_UNKNOWN_FORMAT
1826 //The FFmpegExtractor support all ffmpeg formats!!!
1827 //Unknown format is defined as MEDIA_MIMETYPE_CONTAINER_FFMPEG
1828 const char *container = MEDIA_MIMETYPE_CONTAINER_FFMPEG;
1830 const char *container = NULL;
1833 ALOGV("list the formats suppoted by ffmpeg: ");
1834 ALOGV("========================================");
1835 for (i = 0; i < NELEM(FILE_FORMATS); ++i) {
1836 ALOGV("format_names[%02d]: %s", i, FILE_FORMATS[i].format);
1838 ALOGV("========================================");
1840 for (i = 0; i < NELEM(FILE_FORMATS); ++i) {
1841 int len = strlen(FILE_FORMATS[i].format);
1842 if (!strncasecmp(name, FILE_FORMATS[i].format, len)) {
1843 container = FILE_FORMATS[i].container;
1851 static const char *SniffFFMPEGCommon(const char *url, float *confidence, bool fastMPEG4)
1855 size_t nb_streams = 0;
1856 const char *container = NULL;
1857 AVFormatContext *ic = NULL;
1858 AVDictionary **opts = NULL;
1860 status_t status = initFFmpeg();
1862 ALOGE("could not init ffmpeg");
1866 ic = avformat_alloc_context();
1869 ALOGE("oom for alloc avformat context");
1873 err = avformat_open_input(&ic, url, NULL, NULL);
1876 ALOGE("%s: avformat_open_input failed, err:%s", url, av_err2str(err));
1880 if (ic->iformat != NULL && ic->iformat->name != NULL &&
1881 findMatchingContainer(ic->iformat->name) != NULL &&
1882 !strcasecmp(findMatchingContainer(ic->iformat->name),
1883 MEDIA_MIMETYPE_CONTAINER_MPEG4)) {
1885 container = findMatchingContainer(ic->iformat->name);
1890 opts = setup_find_stream_info_opts(ic, codec_opts);
1891 nb_streams = ic->nb_streams;
1892 err = avformat_find_stream_info(ic, opts);
1894 ALOGE("%s: could not find stream info, err:%s", url, av_err2str(err));
1897 for (i = 0; i < nb_streams; i++) {
1898 av_dict_free(&opts[i]);
1902 av_dump_format(ic, 0, url, 0);
1904 ALOGD("FFmpegExtrator, url: %s, format_name: %s, format_long_name: %s",
1905 url, ic->iformat->name, ic->iformat->long_name);
1907 container = findMatchingContainer(ic->iformat->name);
1909 adjustContainerIfNeeded(&container, ic);
1910 adjustConfidenceIfNeeded(container, ic, confidence);
1915 avformat_close_input(&ic);
1924 static const char *BetterSniffFFMPEG(const sp<DataSource> &source,
1925 float *confidence, sp<AMessage> meta)
1927 const char *ret = NULL;
1928 char url[PATH_MAX] = {0};
1930 ALOGI("android-source:%p", source.get());
1932 // pass the addr of smart pointer("source")
1933 snprintf(url, sizeof(url), "android-source:%p", source.get());
1935 ret = SniffFFMPEGCommon(url, confidence, (source->flags() & DataSource::kIsCachingDataSource));
1937 meta->setString("extended-extractor-url", url);
1943 static const char *LegacySniffFFMPEG(const sp<DataSource> &source,
1944 float *confidence, sp<AMessage> meta)
1946 const char *ret = NULL;
1947 char url[PATH_MAX] = {0};
1949 String8 uri = source->getUri();
1950 if (!uri.string()) {
1954 ALOGV("source url:%s", uri.string());
1956 // pass the addr of smart pointer("source") + file name
1957 snprintf(url, sizeof(url), "android-source:%p|file:%s", source.get(), uri.string());
1959 ret = SniffFFMPEGCommon(url, confidence, false);
1961 meta->setString("extended-extractor-url", url);
1968 const sp<DataSource> &source, String8 *mimeType, float *confidence,
1969 sp<AMessage> *meta) {
1970 ALOGV("SniffFFMPEG");
1972 *meta = new AMessage;
1973 *confidence = 0.08f; // be the last resort, by default
1975 const char *container = BetterSniffFFMPEG(source, confidence, *meta);
1977 ALOGW("sniff through BetterSniffFFMPEG failed, try LegacySniffFFMPEG");
1978 container = LegacySniffFFMPEG(source, confidence, *meta);
1980 ALOGV("sniff through LegacySniffFFMPEG success");
1983 ALOGV("sniff through BetterSniffFFMPEG success");
1986 if (container == NULL) {
1987 ALOGD("SniffFFMPEG failed to sniff this source");
1993 ALOGD("ffmpeg detected media content as '%s' with confidence %.2f",
1994 container, *confidence);
1996 /* use MPEG4Extractor(not extended extractor) for HTTP source only */
1997 if (!strcasecmp(container, MEDIA_MIMETYPE_CONTAINER_MPEG4)
1998 && (source->flags() & DataSource::kIsCachingDataSource)) {
1999 ALOGI("support container: %s, but it is caching data source, "
2000 "Don't use ffmpegextractor", container);
2006 mimeType->setTo(container);
2008 (*meta)->setString("extended-extractor", "extended-extractor");
2009 (*meta)->setString("extended-extractor-subtype", "ffmpegextractor");
2010 (*meta)->setString("extended-extractor-mime", container);
2013 char value[PROPERTY_VALUE_MAX];
2014 property_get("sys.media.parser.ffmpeg", value, "0");
2016 ALOGD("[debug] use ffmpeg parser");
2017 *confidence = 0.88f;
2020 if (*confidence > 0.08f) {
2021 (*meta)->setString("extended-extractor-use", "ffmpegextractor");
2027 MediaExtractor *CreateFFmpegExtractor(const sp<DataSource> &source, const char *mime, const sp<AMessage> &meta) {
2028 MediaExtractor *ret = NULL;
2030 if (meta.get() && meta->findString("extended-extractor", ¬use) && (
2031 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_MPEG) ||
2032 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_AAC) ||
2033 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_VORBIS) ||
2034 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_FLAC) ||
2035 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_AC3) ||
2036 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_APE) ||
2037 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_DTS) ||
2038 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_MPEG_LAYER_II) ||
2039 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_RA) ||
2040 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_WMA) ||
2041 !strcasecmp(mime, MEDIA_MIMETYPE_AUDIO_FFMPEG) ||
2042 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MPEG4) ||
2043 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MOV) ||
2044 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MATROSKA) ||
2045 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_TS) ||
2046 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MPEG2PS) ||
2047 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_AVI) ||
2048 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_ASF) ||
2049 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_WEBM) ||
2050 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_WMV) ||
2051 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MPG) ||
2052 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_FLV) ||
2053 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_DIVX) ||
2054 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_RM) ||
2055 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_WAV) ||
2056 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_FLAC) ||
2057 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_APE) ||
2058 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_DTS) ||
2059 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_MP2) ||
2060 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_RA) ||
2061 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_OGG) ||
2062 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_VC1) ||
2063 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_HEVC) ||
2064 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_WMA) ||
2065 !strcasecmp(mime, MEDIA_MIMETYPE_CONTAINER_FFMPEG))) {
2066 ret = new FFmpegExtractor(source, meta);
2069 ALOGD("%ssupported mime: %s", (ret ? "" : "un"), mime);
2073 } // namespace android
2075 extern "C" void getExtractorPlugin(android::MediaExtractor::Plugin *plugin)
2077 plugin->sniff = android::SniffFFMPEG;
2078 plugin->create = android::CreateFFmpegExtractor;