Prepare the data
# Model / data parameters
num_classes <- 10
input_shape <- c(28, 28, 1)
# Load the data and split it between train and test sets
c(c(x_train, y_train), c(x_test, y_test)) %<-% dataset_mnist()
# Scale images to the [0, 1] range
x_train <- x_train / 255
x_test <- x_test / 255
# Make sure images have shape (28, 28, 1)
x_train <- op_expand_dims(x_train, -1)
x_test <- op_expand_dims(x_test, -1)
dim(x_train)
## [1] 60000 28 28 1
## [1] 10000 28 28 1
# convert class vectors to binary class matrices
y_train <- to_categorical(y_train, num_classes)
y_test <- to_categorical(y_test, num_classes)
Build the model
## Model: "sequential"
## ┏━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━┳━━━━━━━━━━━━━━━━━━━━━━━━┳━━━━━━━━━━━━━━━┓
## ┃ Layer (type) ┃ Output Shape ┃ Param # ┃
## ┡━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━╇━━━━━━━━━━━━━━━━━━━━━━━━╇━━━━━━━━━━━━━━━┩
## │ conv2d_1 (Conv2D) │ (None, 26, 26, 32) │ 320 │
## ├─────────────────────────────────┼────────────────────────┼───────────────┤
## │ max_pooling2d_1 (MaxPooling2D) │ (None, 13, 13, 32) │ 0 │
## ├─────────────────────────────────┼────────────────────────┼───────────────┤
## │ conv2d (Conv2D) │ (None, 11, 11, 64) │ 18,496 │
## ├─────────────────────────────────┼────────────────────────┼───────────────┤
## │ max_pooling2d (MaxPooling2D) │ (None, 5, 5, 64) │ 0 │
## ├─────────────────────────────────┼────────────────────────┼───────────────┤
## │ flatten (Flatten) │ (None, 1600) │ 0 │
## ├─────────────────────────────────┼────────────────────────┼───────────────┤
## │ dropout (Dropout) │ (None, 1600) │ 0 │
## ├─────────────────────────────────┼────────────────────────┼───────────────┤
## │ dense (Dense) │ (None, 10) │ 16,010 │
## └─────────────────────────────────┴────────────────────────┴───────────────┘
## Total params: 34,826 (136.04 KB)
## Trainable params: 34,826 (136.04 KB)
## Non-trainable params: 0 (0.00 B)
Train the model
batch_size <- 128
epochs <- 15
model |> compile(
loss = "categorical_crossentropy",
optimizer = "adam",
metrics = "accuracy"
)
model |> fit(
x_train, y_train,
batch_size = batch_size,
epochs = epochs,
validation_split = 0.1
)
## Epoch 1/15
## 422/422 - 5s - 11ms/step - accuracy: 0.8846 - loss: 0.3816 - val_accuracy: 0.9778 - val_loss: 0.0807
## Epoch 2/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9644 - loss: 0.1151 - val_accuracy: 0.9865 - val_loss: 0.0545
## Epoch 3/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9735 - loss: 0.0839 - val_accuracy: 0.9875 - val_loss: 0.0453
## Epoch 4/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9780 - loss: 0.0691 - val_accuracy: 0.9893 - val_loss: 0.0411
## Epoch 5/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9811 - loss: 0.0609 - val_accuracy: 0.9905 - val_loss: 0.0372
## Epoch 6/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9822 - loss: 0.0562 - val_accuracy: 0.9910 - val_loss: 0.0358
## Epoch 7/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9849 - loss: 0.0490 - val_accuracy: 0.9920 - val_loss: 0.0321
## Epoch 8/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9844 - loss: 0.0484 - val_accuracy: 0.9918 - val_loss: 0.0327
## Epoch 9/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9860 - loss: 0.0440 - val_accuracy: 0.9922 - val_loss: 0.0313
## Epoch 10/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9866 - loss: 0.0405 - val_accuracy: 0.9925 - val_loss: 0.0317
## Epoch 11/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9876 - loss: 0.0389 - val_accuracy: 0.9920 - val_loss: 0.0311
## Epoch 12/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9882 - loss: 0.0370 - val_accuracy: 0.9923 - val_loss: 0.0303
## Epoch 13/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9879 - loss: 0.0362 - val_accuracy: 0.9930 - val_loss: 0.0274
## Epoch 14/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9889 - loss: 0.0334 - val_accuracy: 0.9935 - val_loss: 0.0284
## Epoch 15/15
## 422/422 - 1s - 2ms/step - accuracy: 0.9898 - loss: 0.0309 - val_accuracy: 0.9928 - val_loss: 0.0285
Evaluate the trained model
score <- model |> evaluate(x_test, y_test, verbose = 0)
score
## $accuracy
## [1] 0.9911
##
## $loss
## [1] 0.02562383