In [1]:
from keras.preprocessing.image import ImageDataGenerator
from keras.layers.pooling import AveragePooling2D
from tensorflow.keras.applications import MobileNetV2
from keras.layers.core import Dropout
from keras.layers.core import Flatten
from keras.layers.core import Dense
from keras.layers import Input
from keras.models import Model
from tensorflow.keras.utils import to_categorical
from tensorflow.keras.optimizers import SGD
from tensorflow.keras.utils import to_categorical 
from sklearn.preprocessing import LabelBinarizer
from sklearn.model_selection import train_test_split
from sklearn.metrics import classification_report
from imutils import paths
from tqdm import tqdm
import matplotlib.pyplot as plt
import numpy as np
import warnings
import argparse
import pickle
import cv2
import os
import sys
import glob
import matplotlib
matplotlib.use("Agg")

In [2]:
warnings.filterwarnings('ignore',category=FutureWarning)
warnings.filterwarnings('ignore',category=DeprecationWarning)

In [3]:
args = {
    "dataset": "./SCVD_mini/Frames",
    "model": "./MoBiLSTM5x3/models/MoBiLSTM5x3.h5",
    "label-bin": "./MoBiLSTM5x3/label-bin/MoBiLSTM5x3.pickle",
    "epochs": 30,
    "plot": "./MoBiLSTM5x3/plots/MoBiLSTM5x3.png"
}

In [4]:
# initialize the set of labels from the spots activity dataset we are
# going to train our network on
LABELS = set(["Violence", "NonViolence", "WeaponViolence"])

# grab the list of images in our dataset directory, then initialize
# the list of data (i.e., images) and class images
print('-'*100)
print("[INFO] loading images...")
print('-'*100)
from imutils import paths
imagePaths = list(paths.list_images(args["dataset"]))
print(imagePaths)
data = []
labels = []

----------------------------------------------------------------------------------------------------
[INFO] loading images...
----------------------------------------------------------------------------------------------------
['./SCVD_mini/5x3\\NonViolence\\nv1_merge_0.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_1.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_2.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_3.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_4.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_5.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_6.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_7.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_8.jpg', './SCVD_mini/5x3\\NonViolence\\nv1_merge_9.jpg', './SCVD_mini/5x3\\NonViolence\\nv2_merge_0.jpg', './SCVD_mini/5x3\\NonViolence\\nv2_merge_1.jpg', './SCVD_mini/5x3\\NonViolence\\nv2_merge_2.jpg', './SCVD_mini/5x3\\NonViolence\\nv2_merge_3.jpg', './SCVD_mini/5x3\\NonViolence\\nv2_merge_4.jpg', './SCVD_mini/5x3\\NonViolence\\nv2_me

In [5]:
# loop over the image paths
for imagePath in tqdm(imagePaths[::]):
    # imagePath : file name ex) V_123
    # extract the class label from the filename
    label = imagePath.split(os.path.sep)[-2] # Violence / NonViolence

    # if the label of the current image is not part of of the labels
    # are interested in, then ignore the image
    if label not in LABELS:
        continue

    # load the image, convert it to RGB channel ordering, and resize
    # it to be a fixed 224x224 pixels, ignoring aspect ratio
    image = cv2.imread(imagePath)
    image = cv2.cvtColor(image, cv2.COLOR_BGR2RGB)
    image = cv2.resize(image, (224, 224))

    # update the data and labels lists, respectively
    data.append(image)
    labels.append(label)

100%|██████████████████████████████████████████████████████████████████████████████████| 88/88 [00:56<00:00,  1.57it/s]


In [40]:
# convert the data and labels to NumPy arrays
dt = np.array(data).reshape(224,224,3)
lbls = np.array(labels)


# perform one-hot encoding on the labels
lb = LabelBinarizer()
lbls_fit = lb.fit_transform(lbls)
#lbls_cat = to_categorical(lbls_fit, num_classes=3)

# partition the data into training and testing splits using 70% of
# the data for training and the remaining 30% for testing
(trainX, testX, trainY, testY) = train_test_split(dt, lbls_fit, test_size=0.2, stratify=lbls_fit, random_state=33)

ValueError: cannot reshape array of size 13246464 into shape (10,224,224,3)

In [34]:
# initialize the training data augmentation object
trainAug = ImageDataGenerator(
    #rotation_range=30,
    zoom_range=0.15,
    #width_shift_range=0.2,
    #height_shift_range=0.2,
    #shear_range=0.15,
    horizontal_flip=True,
    fill_mode="nearest")


# initialize the validation/testing data augmentation object (which
# we'll be adding mean subtraction to)
valAug = ImageDataGenerator()

In [35]:
# define the ImageNet mean subtraction (in RGB order) and set the
# the mean subtraction value for each of the data augmentation
# objects
mean = np.array([123.68, 116.779, 103.939], dtype="float32")
trainAug.mean = mean
valAug.mean = mean

In [37]:
from tensorflow.keras.models import Sequential
from tensorflow.keras.layers import *
from tensorflow.keras.utils import plot_model
model = Sequential()
dt=dt.astype('float32')/255

# load the InceptionV3 network, ensuring the head FC layer sets are left
# off
mobilenet = MobileNetV2(
    alpha=1.0,
    include_top=True,
    weights=None,
    input_tensor=None,
    pooling=None,
    classes=3,
    classifier_activation="softmax")

# loop over all layers in the base model and freeze them so they will
# *not* be updated during the training process
mobilenet.trainable = True

model.add(mobilenet,input_shape=dt.shape[1:]))
model.add(Flatten())

model.add(Bidirectional(LSTM(32)))
#model.add(LSTM(32))
model.add(Dense(64,activation='relu'))
model.add(Dense(32,activation='relu'))
#model.add(Flatten())
model.add(Dense(2,activation='sigmoid'))

ValueError: Input 0 of layer "bidirectional_7" is incompatible with the layer: expected ndim=3, found ndim=2. Full shape received: (None, 3)

In [None]:
# compile our model (this needs to be done after our setting our
# layers to being non-trainable)
print('-'*100)
print("[INFO] compiling model...")
print('-'*100)
model.compile(loss='sparse_categorical_crossentropy',optimizer='SGD',metrics=['accuracy'])
print(model.summary())

Object `reshape` not found.
