21#define DEFAULT_INPUT_NAME "transforms.trf"
23#include <vid.stab/libvidstab.h>
47#define OFFSET(x) offsetof(TransformContext, x)
48#define OFFSETC(x) (offsetof(TransformContext, conf)+offsetof(VSTransformConfig, x))
49#define FLAGS AV_OPT_FLAG_FILTERING_PARAM|AV_OPT_FLAG_VIDEO_PARAM
52 {
"input",
"set path to the file storing the transforms",
OFFSET(input),
54 {
"smoothing",
"set number of frames*2 + 1 used for lowpass filtering",
OFFSETC(smoothing),
57 {
"optalgo",
"set camera path optimization algo",
OFFSETC(camPathAlgo),
59 {
"opt",
"global optimization", 0,
61 {
"gauss",
"gaussian kernel", 0,
63 {
"avg",
"simple averaging on motion", 0,
66 {
"maxshift",
"set maximal number of pixels to translate image",
OFFSETC(maxShift),
68 {
"maxangle",
"set maximal angle in rad to rotate image",
OFFSETC(maxAngle),
71 {
"crop",
"set cropping mode",
OFFSETC(crop),
73 {
"keep",
"keep border", 0,
75 {
"black",
"black border", 0,
82 {
"zoom",
"set percentage to zoom (>0: zoom in, <0: zoom out",
OFFSETC(
zoom),
84 {
"optzoom",
"set optimal zoom (0: nothing, 1: optimal static zoom, 2: optimal dynamic zoom)",
OFFSETC(optZoom),
86 {
"zoomspeed",
"for adative zoom: percent to zoom maximally each frame",
OFFSETC(zoomSpeed),
89 {
"interpol",
"set type of interpolation",
OFFSETC(interpolType),
91 {
"no",
"no interpolation", 0,
93 {
"linear",
"linear (horizontal)", 0,
95 {
"bilinear",
"bi-linear", 0,
97 {
"bicubic",
"bi-cubic", 0,
100 {
"tripod",
"enable virtual tripod mode (same as relative=0:smoothing=0)",
OFFSET(tripod),
102 {
"debug",
"enable debug mode and writer global motions information to file",
OFFSET(debug),
113 tc->
class = &vidstabtransform_class;
122 vsTransformDataCleanup(&tc->
td);
123 vsTransformationsCleanup(&tc->
trans);
135 VSTransformData *td = &(tc->
td);
140 if (!vsFrameInfoInit(&fi_src, inlink->
w, inlink->
h,
142 !vsFrameInfoInit(&fi_dest, inlink->
w, inlink->
h,
150 fi_src.log2ChromaW !=
desc->log2_chroma_w ||
151 fi_src.log2ChromaH !=
desc->log2_chroma_h) {
155 fi_src.log2ChromaW,
desc->log2_chroma_w,
156 fi_src.log2ChromaH,
desc->log2_chroma_h);
161 tc->
conf.modName =
"vidstabtransform";
165 tc->
conf.relative = 0;
166 tc->
conf.smoothing = 0;
168 tc->
conf.simpleMotionCalculation = 0;
170 tc->
conf.smoothZoom = 0;
172 if (vsTransformDataInit(td, &tc->
conf, &fi_src, &fi_dest) != VS_OK) {
177 vsTransformGetConfig(&tc->
conf, td);
182 tc->
conf.camPathAlgo == VSOptimalL1 ?
"opt" :
183 (tc->
conf.camPathAlgo == VSGaussian ?
"gauss" :
"avg"));
191 tc->
conf.optZoom == 1 ?
"Static (1)" : (tc->
conf.optZoom == 2 ?
"Dynamic (2)" :
"Off (0)"));
192 if (tc->
conf.optZoom == 2)
202 VSManyLocalMotions mlms;
203 if (vsReadLocalMotionsFile(
f, &mlms) == VS_OK) {
205 if (vsLocalmotions2Transforms(td, &mlms, &tc->
trans) != VS_OK) {
210 if (!vsReadOldTransforms(td,
f, &tc->
trans)) {
218 if (vsPreprocessTransforms(td, &tc->
trans) != VS_OK) {
232 VSTransformData* td = &(tc->
td);
238 for (plane = 0; plane < vsTransformGetSrcFrameInfo(td)->planes; plane++) {
239 inframe.data[plane] = in->
data[plane];
240 inframe.linesize[plane] = in->
linesize[plane];
242 vsTransformPrepare(td, &inframe, &inframe);
244 vsDoTransform(td, vsGetNextTransform(td, &tc->
trans));
246 vsTransformFinish(td);
273 .p.name =
"vidstabtransform",
275 "pass 2 of 2 for stabilization "
276 "(see vidstabdetect for pass 1)."),
277 .p.priv_class = &vidstabtransform_class,
static int config_input(AVFilterLink *inlink)
const FFFilter ff_vf_vidstabtransform
static void invert(float *h, int n)
int ff_filter_frame(AVFilterLink *link, AVFrame *frame)
Send a frame of data to the next filter.
Main libavfilter public API header.
static IPT relative(const CmsCtx *ctx, IPT ipt)
common internal and external API header
static int filter_frame(DBEDecodeContext *s, AVFrame *frame)
int(* init)(AVBSFContext *ctx)
@ AV_OPT_TYPE_CONST
Special option type for declaring named constants.
@ AV_OPT_TYPE_INT
Underlying C type is int.
@ AV_OPT_TYPE_DOUBLE
Underlying C type is double.
@ AV_OPT_TYPE_BOOL
Underlying C type is int.
@ AV_OPT_TYPE_STRING
Underlying C type is a uint8_t* that is either NULL or points to a C string allocated with the av_mal...
#define AV_LOG_VERBOSE
Detailed information.
#define AV_LOG_INFO
Standard information.
#define AV_LOG_ERROR
Something went wrong and cannot losslessly be recovered.
static av_cold void uninit(AVBitStreamFilterContext *ctx)
#define FILTER_INPUTS(array)
#define FILTER_OUTPUTS(array)
#define AVFILTERPAD_FLAG_NEEDS_WRITABLE
The filter expects writable frames from its input link, duplicating data buffers if needed.
#define FILTER_PIXFMTS_ARRAY(array)
#define AVFILTER_DEFINE_CLASS(fname)
FILE * avpriv_fopen_utf8(const char *path, const char *mode)
Open a file using a UTF-8 filename.
#define NULL_IF_CONFIG_SMALL(x)
Return NULL if CONFIG_SMALL is true, otherwise the argument without modification.
int av_get_bits_per_pixel(const AVPixFmtDescriptor *pixdesc)
Return the number of bits per pixel used by the pixel format described by pixdesc.
const AVPixFmtDescriptor * av_pix_fmt_desc_get(enum AVPixelFormat pix_fmt)
#define AV_PIX_FMT_FLAG_PLANAR
At least one pixel component is not in the first data plane.
Describe the class of an AVClass context structure.
A link between two filters.
int w
agreed upon image width
int h
agreed upon image height
AVFilterContext * dst
dest filter
int format
agreed upon media format
A filter pad used for either input or output.
This structure describes decoded (raw) audio or video data.
uint8_t * data[AV_NUM_DATA_POINTERS]
pointer to the picture/channel planes.
int linesize[AV_NUM_DATA_POINTERS]
For video, a positive or negative value, which is typically indicating the size in bytes of each pict...
Descriptor that unambiguously describes how the bits of a pixel are stored in the up to 4 data planes...
static AVFormatContext * ctx
static void zoom(float *u, float *v, float amount)
const AVFilterPad ff_video_default_filterpad[1]
An AVFilterPad array whose only entry has name "default" and is of type AVMEDIA_TYPE_VIDEO.
VSPixelFormat ff_av2vs_pixfmt(AVFilterContext *ctx, enum AVPixelFormat pf)
convert AV's pixelformat to vid.stab pixelformat
enum AVPixelFormat ff_vidstab_pix_fmts[]
void ff_vs_init(void)
sets the memory allocation function and logging constants to av versions